~linaro-toolchain-dev/cortex-strings/trunk

1 by Michael Hope
Pulled in the initial versions
1
/*
2
 * Copyright (c) 2008 ARM Ltd
3
 * All rights reserved.
4
 *
5
 * Redistribution and use in source and binary forms, with or without
6
 * modification, are permitted provided that the following conditions
7
 * are met:
8
 * 1. Redistributions of source code must retain the above copyright
9
 *    notice, this list of conditions and the following disclaimer.
10
 * 2. Redistributions in binary form must reproduce the above copyright
11
 *    notice, this list of conditions and the following disclaimer in the
12
 *    documentation and/or other materials provided with the distribution.
13
 * 3. The name of the company may not be used to endorse or promote
14
 *    products derived from this software without specific prior written
15
 *    permission.
16
 *
17
 * THIS SOFTWARE IS PROVIDED BY ARM LTD ``AS IS'' AND ANY EXPRESS OR IMPLIED
18
 * WARRANTIES, INCLUDING, BUT NOT LIMITED TO, THE IMPLIED WARRANTIES OF
19
 * MERCHANTABILITY AND FITNESS FOR A PARTICULAR PURPOSE ARE DISCLAIMED.
20
 * IN NO EVENT SHALL ARM LTD BE LIABLE FOR ANY DIRECT, INDIRECT, INCIDENTAL,
21
 * SPECIAL, EXEMPLARY, OR CONSEQUENTIAL DAMAGES (INCLUDING, BUT NOT LIMITED
22
 * TO, PROCUREMENT OF SUBSTITUTE GOODS OR SERVICES; LOSS OF USE, DATA, OR
23
 * PROFITS; OR BUSINESS INTERRUPTION) HOWEVER CAUSED AND ON ANY THEORY OF
24
 * LIABILITY, WHETHER IN CONTRACT, STRICT LIABILITY, OR TORT (INCLUDING
25
 * NEGLIGENCE OR OTHERWISE) ARISING IN ANY WAY OUT OF THE USE OF THIS
26
 * SOFTWARE, EVEN IF ADVISED OF THE POSSIBILITY OF SUCH DAMAGE.
27
 */
28
29
#include "arm_asm.h"
30
31
#ifdef __thumb2__
32
#define magic1(REG) "#0x01010101"
33
#define magic2(REG) "#0x80808080"
34
#else
35
#define magic1(REG) #REG
36
#define magic2(REG) #REG ", lsl #7"
37
#endif
38
39
char* __attribute__((naked))
40
strcpy (char* dst, const char* src)
41
{
42
  asm (
43
#if !(defined(__OPTIMIZE_SIZE__) || defined (PREFER_SIZE_OVER_SPEED) || \
44
      (defined (__thumb__) && !defined (__thumb2__)))
45
       "optpld	r1\n\t"
46
       "eor	r2, r0, r1\n\t"
47
       "mov	ip, r0\n\t"
48
       "tst	r2, #3\n\t"
49
       "bne	4f\n\t"
50
       "tst	r1, #3\n\t"
51
       "bne	3f\n"
52
  "5:\n\t"
53
#ifndef __thumb2__
54
       "str	r5, [sp, #-4]!\n\t"
55
       "mov	r5, #0x01\n\t"
56
       "orr	r5, r5, r5, lsl #8\n\t"
57
       "orr	r5, r5, r5, lsl #16\n\t"
58
#endif
59
60
       "str	r4, [sp, #-4]!\n\t"
61
       "tst	r1, #4\n\t"
62
       "ldr	r3, [r1], #4\n\t"
63
       "beq	2f\n\t"
64
       "sub	r2, r3, "magic1(r5)"\n\t"
65
       "bics	r2, r2, r3\n\t"
66
       "tst	r2, "magic2(r5)"\n\t"
67
       "itt	eq\n\t"
68
       "streq	r3, [ip], #4\n\t"
69
       "ldreq	r3, [r1], #4\n"
70
       "bne	1f\n\t"
71
       /* Inner loop.  We now know that r1 is 64-bit aligned, so we
72
	  can safely fetch up to two words.  This allows us to avoid
73
	  load stalls.  */
74
       ".p2align 2\n"
75
  "2:\n\t"
76
       "optpld	r1, #8\n\t"
77
       "ldr	r4, [r1], #4\n\t"
78
       "sub	r2, r3, "magic1(r5)"\n\t"
79
       "bics	r2, r2, r3\n\t"
80
       "tst	r2, "magic2(r5)"\n\t"
81
       "sub	r2, r4, "magic1(r5)"\n\t"
82
       "bne	1f\n\t"
83
       "str	r3, [ip], #4\n\t"
84
       "bics	r2, r2, r4\n\t"
85
       "tst	r2, "magic2(r5)"\n\t"
86
       "itt	eq\n\t"
87
       "ldreq	r3, [r1], #4\n\t"
88
       "streq	r4, [ip], #4\n\t"
89
       "beq	2b\n\t"
90
       "mov	r3, r4\n"
91
  "1:\n\t"
92
#ifdef __ARMEB__
93
       "rors	r3, r3, #24\n\t"
94
#endif
95
       "strb	r3, [ip], #1\n\t"
96
       "tst	r3, #0xff\n\t"
97
#ifdef __ARMEL__
98
       "ror	r3, r3, #8\n\t"
99
#endif
100
       "bne	1b\n\t"
101
       "ldr	r4, [sp], #4\n\t"
102
#ifndef __thumb2__
103
       "ldr	r5, [sp], #4\n\t"
104
#endif
105
       "RETURN\n"
106
107
       /* Strings have the same offset from word alignment, but it's
108
	  not zero.  */
109
  "3:\n\t"
110
       "tst	r1, #1\n\t"
111
       "beq	1f\n\t"
112
       "ldrb	r2, [r1], #1\n\t"
113
       "strb	r2, [ip], #1\n\t"
114
       "cmp	r2, #0\n\t"
115
       "it	eq\n"
116
       "RETURN	eq\n"
117
  "1:\n\t"
118
       "tst	r1, #2\n\t"
119
       "beq	5b\n\t"
120
       "ldrh	r2, [r1], #2\n\t"
121
#ifdef __ARMEB__
122
       "tst	r2, #0xff00\n\t"
123
       "iteet	ne\n\t"
124
       "strneh	r2, [ip], #2\n\t"
125
       "lsreq	r2, r2, #8\n\t"
126
       "streqb	r2, [ip]\n\t"
127
       "tstne	r2, #0xff\n\t"
128
#else
129
       "tst	r2, #0xff\n\t"
130
       "itet	ne\n\t"
131
       "strneh	r2, [ip], #2\n\t"
132
       "streqb	r2, [ip]\n\t"
133
       "tstne	r2, #0xff00\n\t"
134
#endif
135
       "bne	5b\n\t"
136
       "RETURN\n"
137
138
       /* src and dst do not have a common word-alignement.  Fall back to
139
	  byte copying.  */
140
  "4:\n\t"
141
       "ldrb	r2, [r1], #1\n\t"
142
       "strb	r2, [ip], #1\n\t"
143
       "cmp	r2, #0\n\t"
144
       "bne	4b\n\t"
145
       "RETURN"
146
147
#elif !defined (__thumb__) || defined (__thumb2__)
148
       "mov	r3, r0\n\t"
149
  "1:\n\t"
150
       "ldrb	r2, [r1], #1\n\t"
151
       "strb	r2, [r3], #1\n\t"
152
       "cmp	r2, #0\n\t"
153
       "bne	1b\n\t"
154
       "RETURN"
155
#else
156
       "mov	r3, r0\n\t"
157
  "1:\n\t"
158
       "ldrb	r2, [r1]\n\t"
159
       "add	r1, r1, #1\n\t"
160
       "strb	r2, [r3]\n\t"
161
       "add	r3, r3, #1\n\t"
162
       "cmp	r2, #0\n\t"
163
       "bne	1b\n\t"
164
       "RETURN"
165
#endif
166
       );
167
}