annotate src/search.c @ 33881:7a9648a6eafb

*** empty log message ***
author Jason Rumney <jasonr@gnu.org>
date Sat, 25 Nov 2000 16:24:43 +0000
parents 9ec478daa468
children 23a62cf7d0eb
Ignore whitespace changes - Everywhere: Within whitespace: At end of lines:
rev   line source
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1 /* String search routines for GNU Emacs.
26088
b7aa6ac26872 Add support for large files, 64-bit Solaris, system locale codings.
Paul Eggert <eggert@twinsun.com>
parents: 25663
diff changeset
2 Copyright (C) 1985, 86,87,93,94,97,98, 1999 Free Software Foundation, Inc.
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
3
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
4 This file is part of GNU Emacs.
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
5
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
6 GNU Emacs is free software; you can redistribute it and/or modify
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
7 it under the terms of the GNU General Public License as published by
12244
ac7375e60931 Update GPL to version 2.
Karl Heuer <kwzh@gnu.org>
parents: 12148
diff changeset
8 the Free Software Foundation; either version 2, or (at your option)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
9 any later version.
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
10
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
11 GNU Emacs is distributed in the hope that it will be useful,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
12 but WITHOUT ANY WARRANTY; without even the implied warranty of
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
13 MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
14 GNU General Public License for more details.
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
15
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
16 You should have received a copy of the GNU General Public License
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
17 along with GNU Emacs; see the file COPYING. If not, write to
14186
ee40177f6c68 Update FSF's address in the preamble.
Erik Naggum <erik@naggum.no>
parents: 14086
diff changeset
18 the Free Software Foundation, Inc., 59 Temple Place - Suite 330,
ee40177f6c68 Update FSF's address in the preamble.
Erik Naggum <erik@naggum.no>
parents: 14086
diff changeset
19 Boston, MA 02111-1307, USA. */
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
20
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
21
4696
1fc792473491 Include <config.h> instead of "config.h".
Roland McGrath <roland@gnu.org>
parents: 4635
diff changeset
22 #include <config.h>
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
23 #include "lisp.h"
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
24 #include "syntax.h"
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
25 #include "category.h"
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
26 #include "buffer.h"
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
27 #include "charset.h"
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
28 #include "region-cache.h"
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
29 #include "commands.h"
2439
b6c62e4abf59 Put interrupt input blocking in a separate file from xterm.h.
Jim Blandy <jimb@redhat.com>
parents: 2393
diff changeset
30 #include "blockinput.h"
20347
d8e5f3c1618b Include "intervals.h" for prototypes.
Andreas Schwab <schwab@suse.de>
parents: 19541
diff changeset
31 #include "intervals.h"
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
32
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
33 #include <sys/types.h>
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
34 #include "regex.h"
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
35
21457
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
36 #define min(a, b) ((a) < (b) ? (a) : (b))
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
37 #define max(a, b) ((a) > (b) ? (a) : (b))
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
38
16275
a4bcfdc9bb66 (REGEXP_CACHE_SIZE): Increase to 20.
Richard M. Stallman <rms@gnu.org>
parents: 16152
diff changeset
39 #define REGEXP_CACHE_SIZE 20
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
40
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
41 /* If the regexp is non-nil, then the buffer contains the compiled form
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
42 of that regexp, suitable for searching. */
16275
a4bcfdc9bb66 (REGEXP_CACHE_SIZE): Increase to 20.
Richard M. Stallman <rms@gnu.org>
parents: 16152
diff changeset
43 struct regexp_cache
a4bcfdc9bb66 (REGEXP_CACHE_SIZE): Increase to 20.
Richard M. Stallman <rms@gnu.org>
parents: 16152
diff changeset
44 {
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
45 struct regexp_cache *next;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
46 Lisp_Object regexp;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
47 struct re_pattern_buffer buf;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
48 char fastmap[0400];
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
49 /* Nonzero means regexp was compiled to do full POSIX backtracking. */
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
50 char posix;
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
51 };
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
52
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
53 /* The instances of that struct. */
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
54 struct regexp_cache searchbufs[REGEXP_CACHE_SIZE];
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
55
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
56 /* The head of the linked list; points to the most recently used buffer. */
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
57 struct regexp_cache *searchbuf_head;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
58
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
59
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
60 /* Every call to re_match, etc., must pass &search_regs as the regs
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
61 argument unless you can show it is unnecessary (i.e., if re_match
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
62 is certainly going to be called again before region-around-match
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
63 can be called).
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
64
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
65 Since the registers are now dynamically allocated, we need to make
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
66 sure not to refer to the Nth register before checking that it has
708
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
67 been allocated by checking search_regs.num_regs.
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
68
708
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
69 The regex code keeps track of whether it has allocated the search
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
70 buffer using bits in the re_pattern_buffer. This means that whenever
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
71 you compile a new pattern, it completely forgets whether it has
708
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
72 allocated any registers, and will allocate new registers the next
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
73 time you call a searching or matching function. Therefore, we need
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
74 to call re_set_registers after compiling a new pattern or after
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
75 setting the match registers, so that the regex functions will be
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
76 able to free or re-allocate it properly. */
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
77 static struct re_registers search_regs;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
78
727
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
79 /* The buffer in which the last search was performed, or
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
80 Qt if the last search was done in a string;
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
81 Qnil if no searching has been done yet. */
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
82 static Lisp_Object last_thing_searched;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
83
14036
621a575db6f7 Comment fixes.
Karl Heuer <kwzh@gnu.org>
parents: 13295
diff changeset
84 /* error condition signaled when regexp compile_pattern fails */
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
85
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
86 Lisp_Object Qinvalid_regexp;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
87
5556
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
88 static void set_search_regs ();
10055
cb713218845a (save_search_regs): Add declaration.
Richard M. Stallman <rms@gnu.org>
parents: 10032
diff changeset
89 static void save_search_regs ();
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
90 static int simple_search ();
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
91 static int boyer_moore ();
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
92 static int search_buffer ();
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
93
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
94 static void
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
95 matcher_overflow ()
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
96 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
97 error ("Stack overflow in regexp matcher");
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
98 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
99
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
100 /* Compile a regexp and signal a Lisp error if anything goes wrong.
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
101 PATTERN is the pattern to compile.
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
102 CP is the place to put the result.
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
103 TRANSLATE is a translation table for ignoring case, or nil for none.
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
104 REGP is the structure that says where to store the "register"
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
105 values that will result from matching this pattern.
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
106 If it is 0, we should compile the pattern not to record any
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
107 subexpression bounds.
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
108 POSIX is nonzero if we want full backtracking (POSIX style)
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
109 for this pattern. 0 means backtrack only enough to get a valid match.
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
110 MULTIBYTE is nonzero if we want to handle multibyte characters in
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
111 PATTERN. 0 means all multibyte characters are recognized just as
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
112 sequences of binary data. */
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
113
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
114 static void
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
115 compile_pattern_1 (cp, pattern, translate, regp, posix, multibyte)
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
116 struct regexp_cache *cp;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
117 Lisp_Object pattern;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
118 Lisp_Object translate;
708
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
119 struct re_registers *regp;
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
120 int posix;
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
121 int multibyte;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
122 {
21915
8f1159b417c2 (search_buffer): Fix casts when assigning raw_pattern.
Richard M. Stallman <rms@gnu.org>
parents: 21887
diff changeset
123 unsigned char *raw_pattern;
20671
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
124 int raw_pattern_size;
18762
12c0de0113af (compile_pattern_1): Don't declare val with CONST.
Richard M. Stallman <rms@gnu.org>
parents: 18193
diff changeset
125 char *val;
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
126 reg_syntax_t old;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
127
20671
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
128 /* MULTIBYTE says whether the text to be searched is multibyte.
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
129 We must convert PATTERN to match that, or we will not really
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
130 find things right. */
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
131
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
132 if (multibyte == STRING_MULTIBYTE (pattern))
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
133 {
21915
8f1159b417c2 (search_buffer): Fix casts when assigning raw_pattern.
Richard M. Stallman <rms@gnu.org>
parents: 21887
diff changeset
134 raw_pattern = (unsigned char *) XSTRING (pattern)->data;
21244
50929073a0ba Use STRING_BYTES and SET_STRING_BYTES.
Richard M. Stallman <rms@gnu.org>
parents: 21171
diff changeset
135 raw_pattern_size = STRING_BYTES (XSTRING (pattern));
20671
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
136 }
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
137 else if (multibyte)
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
138 {
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
139 raw_pattern_size = count_size_as_multibyte (XSTRING (pattern)->data,
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
140 XSTRING (pattern)->size);
21915
8f1159b417c2 (search_buffer): Fix casts when assigning raw_pattern.
Richard M. Stallman <rms@gnu.org>
parents: 21887
diff changeset
141 raw_pattern = (unsigned char *) alloca (raw_pattern_size + 1);
20671
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
142 copy_text (XSTRING (pattern)->data, raw_pattern,
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
143 XSTRING (pattern)->size, 0, 1);
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
144 }
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
145 else
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
146 {
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
147 /* Converting multibyte to single-byte.
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
148
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
149 ??? Perhaps this conversion should be done in a special way
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
150 by subtracting nonascii-insert-offset from each non-ASCII char,
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
151 so that only the multibyte chars which really correspond to
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
152 the chosen single-byte character set can possibly match. */
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
153 raw_pattern_size = XSTRING (pattern)->size;
21915
8f1159b417c2 (search_buffer): Fix casts when assigning raw_pattern.
Richard M. Stallman <rms@gnu.org>
parents: 21887
diff changeset
154 raw_pattern = (unsigned char *) alloca (raw_pattern_size + 1);
20671
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
155 copy_text (XSTRING (pattern)->data, raw_pattern,
21244
50929073a0ba Use STRING_BYTES and SET_STRING_BYTES.
Richard M. Stallman <rms@gnu.org>
parents: 21171
diff changeset
156 STRING_BYTES (XSTRING (pattern)), 1, 0);
20671
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
157 }
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
158
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
159 cp->regexp = Qnil;
21531
5811a3129878 (compile_pattern, compile_pattern_1): Fix mixing of
Andreas Schwab <schwab@suse.de>
parents: 21514
diff changeset
160 cp->buf.translate = (! NILP (translate) ? translate : make_number (0));
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
161 cp->posix = posix;
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
162 cp->buf.multibyte = multibyte;
2439
b6c62e4abf59 Put interrupt input blocking in a separate file from xterm.h.
Jim Blandy <jimb@redhat.com>
parents: 2393
diff changeset
163 BLOCK_INPUT;
27692
bb0e45f6ca86 * regex.h (RE_SYNTAX_EMACS): Add RE_CHAR_CLASSES and RE_INTERVALS
Stefan Monnier <monnier@iro.umontreal.ca>
parents: 27592
diff changeset
164 old = re_set_syntax (RE_SYNTAX_EMACS
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
165 | (posix ? 0 : RE_NO_POSIX_BACKTRACKING));
21915
8f1159b417c2 (search_buffer): Fix casts when assigning raw_pattern.
Richard M. Stallman <rms@gnu.org>
parents: 21887
diff changeset
166 val = (char *) re_compile_pattern ((char *)raw_pattern,
8f1159b417c2 (search_buffer): Fix casts when assigning raw_pattern.
Richard M. Stallman <rms@gnu.org>
parents: 21887
diff changeset
167 raw_pattern_size, &cp->buf);
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
168 re_set_syntax (old);
2439
b6c62e4abf59 Put interrupt input blocking in a separate file from xterm.h.
Jim Blandy <jimb@redhat.com>
parents: 2393
diff changeset
169 UNBLOCK_INPUT;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
170 if (val)
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
171 Fsignal (Qinvalid_regexp, Fcons (build_string (val), Qnil));
708
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
172
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
173 cp->regexp = Fcopy_sequence (pattern);
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
174 }
708
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
175
22221
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
176 /* Shrink each compiled regexp buffer in the cache
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
177 to the size actually used right now.
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
178 This is called from garbage collection. */
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
179
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
180 void
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
181 shrink_regexp_cache ()
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
182 {
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
183 struct regexp_cache *cp, **cpp;
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
184
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
185 for (cp = searchbuf_head; cp != 0; cp = cp->next)
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
186 {
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
187 cp->buf.allocated = cp->buf.used;
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
188 cp->buf.buffer
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
189 = (unsigned char *) realloc (cp->buf.buffer, cp->buf.used);
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
190 }
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
191 }
239a4b800303 (shrink_regexp_cache): New function.
Richard M. Stallman <rms@gnu.org>
parents: 22082
diff changeset
192
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
193 /* Compile a regexp if necessary, but first check to see if there's one in
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
194 the cache.
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
195 PATTERN is the pattern to compile.
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
196 TRANSLATE is a translation table for ignoring case, or nil for none.
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
197 REGP is the structure that says where to store the "register"
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
198 values that will result from matching this pattern.
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
199 If it is 0, we should compile the pattern not to record any
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
200 subexpression bounds.
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
201 POSIX is nonzero if we want full backtracking (POSIX style)
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
202 for this pattern. 0 means backtrack only enough to get a valid match. */
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
203
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
204 struct re_pattern_buffer *
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
205 compile_pattern (pattern, regp, translate, posix, multibyte)
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
206 Lisp_Object pattern;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
207 struct re_registers *regp;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
208 Lisp_Object translate;
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
209 int posix, multibyte;
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
210 {
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
211 struct regexp_cache *cp, **cpp;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
212
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
213 for (cpp = &searchbuf_head; ; cpp = &cp->next)
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
214 {
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
215 cp = *cpp;
27592
5cd59d1800ad * search.c (compile_pattern): If a cache entry has a nil regexp, fill in that
Ken Raeburn <raeburn@raeburn.org>
parents: 26985
diff changeset
216 /* Entries are initialized to nil, and may be set to nil by
5cd59d1800ad * search.c (compile_pattern): If a cache entry has a nil regexp, fill in that
Ken Raeburn <raeburn@raeburn.org>
parents: 26985
diff changeset
217 compile_pattern_1 if the pattern isn't valid. Don't apply
5cd59d1800ad * search.c (compile_pattern): If a cache entry has a nil regexp, fill in that
Ken Raeburn <raeburn@raeburn.org>
parents: 26985
diff changeset
218 XSTRING in those cases. However, compile_pattern_1 is only
5cd59d1800ad * search.c (compile_pattern): If a cache entry has a nil regexp, fill in that
Ken Raeburn <raeburn@raeburn.org>
parents: 26985
diff changeset
219 applied to the cache entry we pick here to reuse. So nil
5cd59d1800ad * search.c (compile_pattern): If a cache entry has a nil regexp, fill in that
Ken Raeburn <raeburn@raeburn.org>
parents: 26985
diff changeset
220 should never appear before a non-nil entry. */
28507
b6f06a755c7d make_number/XINT/XUINT conversions; EQ/== fixes; ==Qnil -> NILP
Ken Raeburn <raeburn@raeburn.org>
parents: 28387
diff changeset
221 if (NILP (cp->regexp))
27592
5cd59d1800ad * search.c (compile_pattern): If a cache entry has a nil regexp, fill in that
Ken Raeburn <raeburn@raeburn.org>
parents: 26985
diff changeset
222 goto compile_it;
16275
a4bcfdc9bb66 (REGEXP_CACHE_SIZE): Increase to 20.
Richard M. Stallman <rms@gnu.org>
parents: 16152
diff changeset
223 if (XSTRING (cp->regexp)->size == XSTRING (pattern)->size
31486
a3dc5f987e8f (compile_pattern): Check the multibyteness of cached
Kenichi Handa <handa@m17n.org>
parents: 29335
diff changeset
224 && STRING_MULTIBYTE (cp->regexp) == STRING_MULTIBYTE (pattern)
16275
a4bcfdc9bb66 (REGEXP_CACHE_SIZE): Increase to 20.
Richard M. Stallman <rms@gnu.org>
parents: 16152
diff changeset
225 && !NILP (Fstring_equal (cp->regexp, pattern))
21531
5811a3129878 (compile_pattern, compile_pattern_1): Fix mixing of
Andreas Schwab <schwab@suse.de>
parents: 21514
diff changeset
226 && EQ (cp->buf.translate, (! NILP (translate) ? translate : make_number (0)))
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
227 && cp->posix == posix
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
228 && cp->buf.multibyte == multibyte)
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
229 break;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
230
27592
5cd59d1800ad * search.c (compile_pattern): If a cache entry has a nil regexp, fill in that
Ken Raeburn <raeburn@raeburn.org>
parents: 26985
diff changeset
231 /* If we're at the end of the cache, compile into the nil cell
5cd59d1800ad * search.c (compile_pattern): If a cache entry has a nil regexp, fill in that
Ken Raeburn <raeburn@raeburn.org>
parents: 26985
diff changeset
232 we found, or the last (least recently used) cell with a
5cd59d1800ad * search.c (compile_pattern): If a cache entry has a nil regexp, fill in that
Ken Raeburn <raeburn@raeburn.org>
parents: 26985
diff changeset
233 string value. */
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
234 if (cp->next == 0)
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
235 {
27592
5cd59d1800ad * search.c (compile_pattern): If a cache entry has a nil regexp, fill in that
Ken Raeburn <raeburn@raeburn.org>
parents: 26985
diff changeset
236 compile_it:
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
237 compile_pattern_1 (cp, pattern, translate, regp, posix, multibyte);
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
238 break;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
239 }
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
240 }
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
241
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
242 /* When we get here, cp (aka *cpp) contains the compiled pattern,
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
243 either because we found it in the cache or because we just compiled it.
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
244 Move it to the front of the queue to mark it as most recently used. */
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
245 *cpp = cp->next;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
246 cp->next = searchbuf_head;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
247 searchbuf_head = cp;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
248
10141
afe81fd385eb (compile_pattern): Call re_set_registers here.
Richard M. Stallman <rms@gnu.org>
parents: 10128
diff changeset
249 /* Advise the searching functions about the space we have allocated
afe81fd385eb (compile_pattern): Call re_set_registers here.
Richard M. Stallman <rms@gnu.org>
parents: 10128
diff changeset
250 for register data. */
afe81fd385eb (compile_pattern): Call re_set_registers here.
Richard M. Stallman <rms@gnu.org>
parents: 10128
diff changeset
251 if (regp)
afe81fd385eb (compile_pattern): Call re_set_registers here.
Richard M. Stallman <rms@gnu.org>
parents: 10128
diff changeset
252 re_set_registers (&cp->buf, regp, regp->num_regs, regp->start, regp->end);
afe81fd385eb (compile_pattern): Call re_set_registers here.
Richard M. Stallman <rms@gnu.org>
parents: 10128
diff changeset
253
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
254 return &cp->buf;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
255 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
256
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
257 /* Error condition used for failing searches */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
258 Lisp_Object Qsearch_failed;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
259
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
260 Lisp_Object
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
261 signal_failure (arg)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
262 Lisp_Object arg;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
263 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
264 Fsignal (Qsearch_failed, Fcons (arg, Qnil));
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
265 return Qnil;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
266 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
267
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
268 static Lisp_Object
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
269 looking_at_1 (string, posix)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
270 Lisp_Object string;
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
271 int posix;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
272 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
273 Lisp_Object val;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
274 unsigned char *p1, *p2;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
275 int s1, s2;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
276 register int i;
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
277 struct re_pattern_buffer *bufp;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
278
10032
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
279 if (running_asynch_code)
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
280 save_search_regs ();
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
281
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
282 CHECK_STRING (string, 0);
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
283 bufp = compile_pattern (string, &search_regs,
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
284 (!NILP (current_buffer->case_fold_search)
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
285 ? DOWNCASE_TABLE : Qnil),
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
286 posix,
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
287 !NILP (current_buffer->enable_multibyte_characters));
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
288
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
289 immediate_quit = 1;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
290 QUIT; /* Do a pending quit right away, to avoid paradoxical behavior */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
291
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
292 /* Get pointers and sizes of the two strings
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
293 that make up the visible portion of the buffer. */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
294
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
295 p1 = BEGV_ADDR;
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
296 s1 = GPT_BYTE - BEGV_BYTE;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
297 p2 = GAP_END_ADDR;
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
298 s2 = ZV_BYTE - GPT_BYTE;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
299 if (s1 < 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
300 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
301 p2 = p1;
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
302 s2 = ZV_BYTE - BEGV_BYTE;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
303 s1 = 0;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
304 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
305 if (s2 < 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
306 {
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
307 s1 = ZV_BYTE - BEGV_BYTE;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
308 s2 = 0;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
309 }
17463
bb9ae80d22e2 (looking_at_1): Set re_match_object.
Richard M. Stallman <rms@gnu.org>
parents: 17284
diff changeset
310
bb9ae80d22e2 (looking_at_1): Set re_match_object.
Richard M. Stallman <rms@gnu.org>
parents: 17284
diff changeset
311 re_match_object = Qnil;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
312
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
313 i = re_match_2 (bufp, (char *) p1, s1, (char *) p2, s2,
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
314 PT_BYTE - BEGV_BYTE, &search_regs,
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
315 ZV_BYTE - BEGV_BYTE);
26985
1121a5da20a5 (looking_at_1): Reset immediate_quit before modifying
Gerd Moellmann <gerd@gnu.org>
parents: 26982
diff changeset
316 immediate_quit = 0;
1121a5da20a5 (looking_at_1): Reset immediate_quit before modifying
Gerd Moellmann <gerd@gnu.org>
parents: 26982
diff changeset
317
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
318 if (i == -2)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
319 matcher_overflow ();
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
320
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
321 val = (0 <= i ? Qt : Qnil);
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
322 if (i >= 0)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
323 for (i = 0; i < search_regs.num_regs; i++)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
324 if (search_regs.start[i] >= 0)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
325 {
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
326 search_regs.start[i]
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
327 = BYTE_TO_CHAR (search_regs.start[i] + BEGV_BYTE);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
328 search_regs.end[i]
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
329 = BYTE_TO_CHAR (search_regs.end[i] + BEGV_BYTE);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
330 }
9278
f2138d548313 (Flooking_at, skip_chars, search_buffer, set_search_regs, Fstore_match_data):
Karl Heuer <kwzh@gnu.org>
parents: 9113
diff changeset
331 XSETBUFFER (last_thing_searched, current_buffer);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
332 return val;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
333 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
334
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
335 DEFUN ("looking-at", Flooking_at, Slooking_at, 1, 1, 0,
11213
d0811ba886f8 (Flooking_at, Fposix_looking_at): Change arg name.
Richard M. Stallman <rms@gnu.org>
parents: 10250
diff changeset
336 "Return t if text after point matches regular expression REGEXP.\n\
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
337 This function modifies the match data that `match-beginning',\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
338 `match-end' and `match-data' access; save and restore the match\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
339 data if you want to preserve them.")
11213
d0811ba886f8 (Flooking_at, Fposix_looking_at): Change arg name.
Richard M. Stallman <rms@gnu.org>
parents: 10250
diff changeset
340 (regexp)
d0811ba886f8 (Flooking_at, Fposix_looking_at): Change arg name.
Richard M. Stallman <rms@gnu.org>
parents: 10250
diff changeset
341 Lisp_Object regexp;
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
342 {
11213
d0811ba886f8 (Flooking_at, Fposix_looking_at): Change arg name.
Richard M. Stallman <rms@gnu.org>
parents: 10250
diff changeset
343 return looking_at_1 (regexp, 0);
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
344 }
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
345
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
346 DEFUN ("posix-looking-at", Fposix_looking_at, Sposix_looking_at, 1, 1, 0,
11213
d0811ba886f8 (Flooking_at, Fposix_looking_at): Change arg name.
Richard M. Stallman <rms@gnu.org>
parents: 10250
diff changeset
347 "Return t if text after point matches regular expression REGEXP.\n\
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
348 Find the longest match, in accord with Posix regular expression rules.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
349 This function modifies the match data that `match-beginning',\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
350 `match-end' and `match-data' access; save and restore the match\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
351 data if you want to preserve them.")
11213
d0811ba886f8 (Flooking_at, Fposix_looking_at): Change arg name.
Richard M. Stallman <rms@gnu.org>
parents: 10250
diff changeset
352 (regexp)
d0811ba886f8 (Flooking_at, Fposix_looking_at): Change arg name.
Richard M. Stallman <rms@gnu.org>
parents: 10250
diff changeset
353 Lisp_Object regexp;
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
354 {
11213
d0811ba886f8 (Flooking_at, Fposix_looking_at): Change arg name.
Richard M. Stallman <rms@gnu.org>
parents: 10250
diff changeset
355 return looking_at_1 (regexp, 1);
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
356 }
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
357
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
358 static Lisp_Object
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
359 string_match_1 (regexp, string, start, posix)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
360 Lisp_Object regexp, string, start;
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
361 int posix;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
362 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
363 int val;
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
364 struct re_pattern_buffer *bufp;
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
365 int pos, pos_byte;
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
366 int i;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
367
10032
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
368 if (running_asynch_code)
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
369 save_search_regs ();
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
370
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
371 CHECK_STRING (regexp, 0);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
372 CHECK_STRING (string, 1);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
373
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
374 if (NILP (start))
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
375 pos = 0, pos_byte = 0;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
376 else
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
377 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
378 int len = XSTRING (string)->size;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
379
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
380 CHECK_NUMBER (start, 2);
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
381 pos = XINT (start);
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
382 if (pos < 0 && -pos <= len)
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
383 pos = len + pos;
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
384 else if (0 > pos || pos > len)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
385 args_out_of_range (string, start);
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
386 pos_byte = string_char_to_byte (string, pos);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
387 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
388
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
389 bufp = compile_pattern (regexp, &search_regs,
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
390 (!NILP (current_buffer->case_fold_search)
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
391 ? DOWNCASE_TABLE : Qnil),
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
392 posix,
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
393 STRING_MULTIBYTE (string));
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
394 immediate_quit = 1;
17463
bb9ae80d22e2 (looking_at_1): Set re_match_object.
Richard M. Stallman <rms@gnu.org>
parents: 17284
diff changeset
395 re_match_object = string;
bb9ae80d22e2 (looking_at_1): Set re_match_object.
Richard M. Stallman <rms@gnu.org>
parents: 17284
diff changeset
396
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
397 val = re_search (bufp, (char *) XSTRING (string)->data,
21244
50929073a0ba Use STRING_BYTES and SET_STRING_BYTES.
Richard M. Stallman <rms@gnu.org>
parents: 21171
diff changeset
398 STRING_BYTES (XSTRING (string)), pos_byte,
50929073a0ba Use STRING_BYTES and SET_STRING_BYTES.
Richard M. Stallman <rms@gnu.org>
parents: 21171
diff changeset
399 STRING_BYTES (XSTRING (string)) - pos_byte,
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
400 &search_regs);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
401 immediate_quit = 0;
727
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
402 last_thing_searched = Qt;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
403 if (val == -2)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
404 matcher_overflow ();
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
405 if (val < 0) return Qnil;
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
406
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
407 for (i = 0; i < search_regs.num_regs; i++)
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
408 if (search_regs.start[i] >= 0)
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
409 {
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
410 search_regs.start[i]
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
411 = string_byte_to_char (string, search_regs.start[i]);
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
412 search_regs.end[i]
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
413 = string_byte_to_char (string, search_regs.end[i]);
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
414 }
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
415
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
416 return make_number (string_byte_to_char (string, val));
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
417 }
842
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
418
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
419 DEFUN ("string-match", Fstring_match, Sstring_match, 2, 3, 0,
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
420 "Return index of start of first match for REGEXP in STRING, or nil.\n\
24433
118e66d79d64 (Fstring_match, Fposix_string_match): Doc fix.
Karl Heuer <kwzh@gnu.org>
parents: 24014
diff changeset
421 Case is ignored if `case-fold-search' is non-nil in the current buffer.\n\
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
422 If third arg START is non-nil, start search at that index in STRING.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
423 For index of first char beyond the match, do (match-end 0).\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
424 `match-end' and `match-beginning' also give indices of substrings\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
425 matched by parenthesis constructs in the pattern.")
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
426 (regexp, string, start)
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
427 Lisp_Object regexp, string, start;
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
428 {
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
429 return string_match_1 (regexp, string, start, 0);
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
430 }
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
431
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
432 DEFUN ("posix-string-match", Fposix_string_match, Sposix_string_match, 2, 3, 0,
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
433 "Return index of start of first match for REGEXP in STRING, or nil.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
434 Find the longest match, in accord with Posix regular expression rules.\n\
24433
118e66d79d64 (Fstring_match, Fposix_string_match): Doc fix.
Karl Heuer <kwzh@gnu.org>
parents: 24014
diff changeset
435 Case is ignored if `case-fold-search' is non-nil in the current buffer.\n\
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
436 If third arg START is non-nil, start search at that index in STRING.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
437 For index of first char beyond the match, do (match-end 0).\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
438 `match-end' and `match-beginning' also give indices of substrings\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
439 matched by parenthesis constructs in the pattern.")
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
440 (regexp, string, start)
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
441 Lisp_Object regexp, string, start;
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
442 {
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
443 return string_match_1 (regexp, string, start, 1);
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
444 }
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
445
842
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
446 /* Match REGEXP against STRING, searching all of STRING,
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
447 and return the index of the match, or negative on failure.
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
448 This does not clobber the match data. */
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
449
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
450 int
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
451 fast_string_match (regexp, string)
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
452 Lisp_Object regexp, string;
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
453 {
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
454 int val;
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
455 struct re_pattern_buffer *bufp;
842
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
456
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
457 bufp = compile_pattern (regexp, 0, Qnil,
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
458 0, STRING_MULTIBYTE (string));
842
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
459 immediate_quit = 1;
17463
bb9ae80d22e2 (looking_at_1): Set re_match_object.
Richard M. Stallman <rms@gnu.org>
parents: 17284
diff changeset
460 re_match_object = string;
bb9ae80d22e2 (looking_at_1): Set re_match_object.
Richard M. Stallman <rms@gnu.org>
parents: 17284
diff changeset
461
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
462 val = re_search (bufp, (char *) XSTRING (string)->data,
21244
50929073a0ba Use STRING_BYTES and SET_STRING_BYTES.
Richard M. Stallman <rms@gnu.org>
parents: 21171
diff changeset
463 STRING_BYTES (XSTRING (string)), 0,
50929073a0ba Use STRING_BYTES and SET_STRING_BYTES.
Richard M. Stallman <rms@gnu.org>
parents: 21171
diff changeset
464 STRING_BYTES (XSTRING (string)), 0);
842
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
465 immediate_quit = 0;
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
466 return val;
5ce0a9ac1ea7 entered into RCS
Richard M. Stallman <rms@gnu.org>
parents: 808
diff changeset
467 }
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
468
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
469 /* Match REGEXP against STRING, searching all of STRING ignoring case,
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
470 and return the index of the match, or negative on failure.
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
471 This does not clobber the match data.
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
472 We assume that STRING contains single-byte characters. */
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
473
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
474 extern Lisp_Object Vascii_downcase_table;
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
475
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
476 int
18193
4e4c8edb56da (fast_c_string_match_ignore_case):
Richard M. Stallman <rms@gnu.org>
parents: 18124
diff changeset
477 fast_c_string_match_ignore_case (regexp, string)
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
478 Lisp_Object regexp;
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
479 char *string;
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
480 {
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
481 int val;
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
482 struct re_pattern_buffer *bufp;
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
483 int len = strlen (string);
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
484
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
485 regexp = string_make_unibyte (regexp);
18193
4e4c8edb56da (fast_c_string_match_ignore_case):
Richard M. Stallman <rms@gnu.org>
parents: 18124
diff changeset
486 re_match_object = Qt;
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
487 bufp = compile_pattern (regexp, 0,
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
488 Vascii_downcase_table, 0,
20671
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
489 0);
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
490 immediate_quit = 1;
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
491 val = re_search (bufp, string, len, 0, len, 0);
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
492 immediate_quit = 0;
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
493 return val;
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
494 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
495
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
496 /* The newline cache: remembering which sections of text have no newlines. */
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
497
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
498 /* If the user has requested newline caching, make sure it's on.
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
499 Otherwise, make sure it's off.
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
500 This is our cheezy way of associating an action with the change of
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
501 state of a buffer-local variable. */
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
502 static void
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
503 newline_cache_on_off (buf)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
504 struct buffer *buf;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
505 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
506 if (NILP (buf->cache_long_line_scans))
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
507 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
508 /* It should be off. */
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
509 if (buf->newline_cache)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
510 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
511 free_region_cache (buf->newline_cache);
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
512 buf->newline_cache = 0;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
513 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
514 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
515 else
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
516 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
517 /* It should be on. */
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
518 if (buf->newline_cache == 0)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
519 buf->newline_cache = new_region_cache ();
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
520 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
521 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
522
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
523
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
524 /* Search for COUNT instances of the character TARGET between START and END.
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
525
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
526 If COUNT is positive, search forwards; END must be >= START.
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
527 If COUNT is negative, search backwards for the -COUNTth instance;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
528 END must be <= START.
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
529 If COUNT is zero, do anything you please; run rogue, for all I care.
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
530
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
531 If END is zero, use BEGV or ZV instead, as appropriate for the
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
532 direction indicated by COUNT.
648
70b112526394 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 639
diff changeset
533
70b112526394 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 639
diff changeset
534 If we find COUNT instances, set *SHORTAGE to zero, and return the
1413
527af9fa8676 Comment fix.
Richard M. Stallman <rms@gnu.org>
parents: 842
diff changeset
535 position after the COUNTth match. Note that for reverse motion
527af9fa8676 Comment fix.
Richard M. Stallman <rms@gnu.org>
parents: 842
diff changeset
536 this is not the same as the usual convention for Emacs motion commands.
648
70b112526394 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 639
diff changeset
537
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
538 If we don't find COUNT instances before reaching END, set *SHORTAGE
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
539 to the number of TARGETs left unfound, and return END.
648
70b112526394 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 639
diff changeset
540
5756
a54c236b43c6 (scan_buffer): New arg ALLOW_QUIT.
Richard M. Stallman <rms@gnu.org>
parents: 5556
diff changeset
541 If ALLOW_QUIT is non-zero, set immediate_quit. That's good to do
a54c236b43c6 (scan_buffer): New arg ALLOW_QUIT.
Richard M. Stallman <rms@gnu.org>
parents: 5556
diff changeset
542 except when inside redisplay. */
a54c236b43c6 (scan_buffer): New arg ALLOW_QUIT.
Richard M. Stallman <rms@gnu.org>
parents: 5556
diff changeset
543
21514
fa9ff387d260 Fix -Wimplicit warnings.
Andreas Schwab <schwab@suse.de>
parents: 21457
diff changeset
544 int
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
545 scan_buffer (target, start, end, count, shortage, allow_quit)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
546 register int target;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
547 int start, end;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
548 int count;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
549 int *shortage;
5756
a54c236b43c6 (scan_buffer): New arg ALLOW_QUIT.
Richard M. Stallman <rms@gnu.org>
parents: 5556
diff changeset
550 int allow_quit;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
551 {
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
552 struct region_cache *newline_cache;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
553 int direction;
648
70b112526394 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 639
diff changeset
554
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
555 if (count > 0)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
556 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
557 direction = 1;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
558 if (! end) end = ZV;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
559 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
560 else
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
561 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
562 direction = -1;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
563 if (! end) end = BEGV;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
564 }
648
70b112526394 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 639
diff changeset
565
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
566 newline_cache_on_off (current_buffer);
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
567 newline_cache = current_buffer->newline_cache;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
568
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
569 if (shortage != 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
570 *shortage = 0;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
571
5756
a54c236b43c6 (scan_buffer): New arg ALLOW_QUIT.
Richard M. Stallman <rms@gnu.org>
parents: 5556
diff changeset
572 immediate_quit = allow_quit;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
573
648
70b112526394 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 639
diff changeset
574 if (count > 0)
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
575 while (start != end)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
576 {
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
577 /* Our innermost scanning loop is very simple; it doesn't know
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
578 about gaps, buffer ends, or the newline cache. ceiling is
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
579 the position of the last character before the next such
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
580 obstacle --- the last character the dumb search loop should
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
581 examine. */
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
582 int ceiling_byte = CHAR_TO_BYTE (end) - 1;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
583 int start_byte = CHAR_TO_BYTE (start);
21457
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
584 int tem;
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
585
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
586 /* If we're looking for a newline, consult the newline cache
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
587 to see where we can avoid some scanning. */
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
588 if (target == '\n' && newline_cache)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
589 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
590 int next_change;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
591 immediate_quit = 0;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
592 while (region_cache_forward
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
593 (current_buffer, newline_cache, start_byte, &next_change))
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
594 start_byte = next_change;
9452
76f75b9091f1 (scan_buffer): After temporarily turning immediate_quit off, turn it
Jim Blandy <jimb@redhat.com>
parents: 9410
diff changeset
595 immediate_quit = allow_quit;
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
596
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
597 /* START should never be after END. */
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
598 if (start_byte > ceiling_byte)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
599 start_byte = ceiling_byte;
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
600
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
601 /* Now the text after start is an unknown region, and
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
602 next_change is the position of the next known region. */
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
603 ceiling_byte = min (next_change - 1, ceiling_byte);
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
604 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
605
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
606 /* The dumb loop can only scan text stored in contiguous
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
607 bytes. BUFFER_CEILING_OF returns the last character
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
608 position that is contiguous, so the ceiling is the
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
609 position after that. */
21457
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
610 tem = BUFFER_CEILING_OF (start_byte);
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
611 ceiling_byte = min (tem, ceiling_byte);
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
612
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
613 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
614 /* The termination address of the dumb loop. */
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
615 register unsigned char *ceiling_addr
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
616 = BYTE_POS_ADDR (ceiling_byte) + 1;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
617 register unsigned char *cursor
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
618 = BYTE_POS_ADDR (start_byte);
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
619 unsigned char *base = cursor;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
620
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
621 while (cursor < ceiling_addr)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
622 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
623 unsigned char *scan_start = cursor;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
624
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
625 /* The dumb loop. */
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
626 while (*cursor != target && ++cursor < ceiling_addr)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
627 ;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
628
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
629 /* If we're looking for newlines, cache the fact that
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
630 the region from start to cursor is free of them. */
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
631 if (target == '\n' && newline_cache)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
632 know_region_cache (current_buffer, newline_cache,
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
633 start_byte + scan_start - base,
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
634 start_byte + cursor - base);
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
635
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
636 /* Did we find the target character? */
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
637 if (cursor < ceiling_addr)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
638 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
639 if (--count == 0)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
640 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
641 immediate_quit = 0;
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
642 return BYTE_TO_CHAR (start_byte + cursor - base + 1);
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
643 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
644 cursor++;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
645 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
646 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
647
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
648 start = BYTE_TO_CHAR (start_byte + cursor - base);
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
649 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
650 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
651 else
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
652 while (start > end)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
653 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
654 /* The last character to check before the next obstacle. */
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
655 int ceiling_byte = CHAR_TO_BYTE (end);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
656 int start_byte = CHAR_TO_BYTE (start);
21457
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
657 int tem;
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
658
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
659 /* Consult the newline cache, if appropriate. */
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
660 if (target == '\n' && newline_cache)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
661 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
662 int next_change;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
663 immediate_quit = 0;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
664 while (region_cache_backward
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
665 (current_buffer, newline_cache, start_byte, &next_change))
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
666 start_byte = next_change;
9452
76f75b9091f1 (scan_buffer): After temporarily turning immediate_quit off, turn it
Jim Blandy <jimb@redhat.com>
parents: 9410
diff changeset
667 immediate_quit = allow_quit;
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
668
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
669 /* Start should never be at or before end. */
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
670 if (start_byte <= ceiling_byte)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
671 start_byte = ceiling_byte + 1;
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
672
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
673 /* Now the text before start is an unknown region, and
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
674 next_change is the position of the next known region. */
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
675 ceiling_byte = max (next_change, ceiling_byte);
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
676 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
677
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
678 /* Stop scanning before the gap. */
21457
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
679 tem = BUFFER_FLOOR_OF (start_byte - 1);
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
680 ceiling_byte = max (tem, ceiling_byte);
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
681
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
682 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
683 /* The termination address of the dumb loop. */
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
684 register unsigned char *ceiling_addr = BYTE_POS_ADDR (ceiling_byte);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
685 register unsigned char *cursor = BYTE_POS_ADDR (start_byte - 1);
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
686 unsigned char *base = cursor;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
687
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
688 while (cursor >= ceiling_addr)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
689 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
690 unsigned char *scan_start = cursor;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
691
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
692 while (*cursor != target && --cursor >= ceiling_addr)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
693 ;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
694
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
695 /* If we're looking for newlines, cache the fact that
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
696 the region from after the cursor to start is free of them. */
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
697 if (target == '\n' && newline_cache)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
698 know_region_cache (current_buffer, newline_cache,
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
699 start_byte + cursor - base,
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
700 start_byte + scan_start - base);
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
701
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
702 /* Did we find the target character? */
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
703 if (cursor >= ceiling_addr)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
704 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
705 if (++count >= 0)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
706 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
707 immediate_quit = 0;
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
708 return BYTE_TO_CHAR (start_byte + cursor - base);
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
709 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
710 cursor--;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
711 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
712 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
713
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
714 start = BYTE_TO_CHAR (start_byte + cursor - base);
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
715 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
716 }
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
717
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
718 immediate_quit = 0;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
719 if (shortage != 0)
648
70b112526394 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 639
diff changeset
720 *shortage = count * direction;
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
721 return start;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
722 }
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
723
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
724 /* Search for COUNT instances of a line boundary, which means either a
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
725 newline or (if selective display enabled) a carriage return.
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
726 Start at START. If COUNT is negative, search backwards.
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
727
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
728 We report the resulting position by calling TEMP_SET_PT_BOTH.
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
729
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
730 If we find COUNT instances. we position after (always after,
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
731 even if scanning backwards) the COUNTth match, and return 0.
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
732
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
733 If we don't find COUNT instances before reaching the end of the
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
734 buffer (or the beginning, if scanning backwards), we return
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
735 the number of line boundaries left unfound, and position at
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
736 the limit we bumped up against.
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
737
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
738 If ALLOW_QUIT is non-zero, set immediate_quit. That's good to do
20547
07053199a368 (scan_newline): Always restore prev value of immediate_quit.
Richard M. Stallman <rms@gnu.org>
parents: 20545
diff changeset
739 except in special cases. */
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
740
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
741 int
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
742 scan_newline (start, start_byte, limit, limit_byte, count, allow_quit)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
743 int start, start_byte;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
744 int limit, limit_byte;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
745 register int count;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
746 int allow_quit;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
747 {
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
748 int direction = ((count > 0) ? 1 : -1);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
749
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
750 register unsigned char *cursor;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
751 unsigned char *base;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
752
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
753 register int ceiling;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
754 register unsigned char *ceiling_addr;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
755
20547
07053199a368 (scan_newline): Always restore prev value of immediate_quit.
Richard M. Stallman <rms@gnu.org>
parents: 20545
diff changeset
756 int old_immediate_quit = immediate_quit;
07053199a368 (scan_newline): Always restore prev value of immediate_quit.
Richard M. Stallman <rms@gnu.org>
parents: 20545
diff changeset
757
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
758 /* If we are not in selective display mode,
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
759 check only for newlines. */
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
760 int selective_display = (!NILP (current_buffer->selective_display)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
761 && !INTEGERP (current_buffer->selective_display));
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
762
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
763 /* The code that follows is like scan_buffer
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
764 but checks for either newline or carriage return. */
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
765
20547
07053199a368 (scan_newline): Always restore prev value of immediate_quit.
Richard M. Stallman <rms@gnu.org>
parents: 20545
diff changeset
766 if (allow_quit)
07053199a368 (scan_newline): Always restore prev value of immediate_quit.
Richard M. Stallman <rms@gnu.org>
parents: 20545
diff changeset
767 immediate_quit++;
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
768
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
769 start_byte = CHAR_TO_BYTE (start);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
770
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
771 if (count > 0)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
772 {
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
773 while (start_byte < limit_byte)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
774 {
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
775 ceiling = BUFFER_CEILING_OF (start_byte);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
776 ceiling = min (limit_byte - 1, ceiling);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
777 ceiling_addr = BYTE_POS_ADDR (ceiling) + 1;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
778 base = (cursor = BYTE_POS_ADDR (start_byte));
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
779 while (1)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
780 {
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
781 while (*cursor != '\n' && ++cursor != ceiling_addr)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
782 ;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
783
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
784 if (cursor != ceiling_addr)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
785 {
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
786 if (--count == 0)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
787 {
20547
07053199a368 (scan_newline): Always restore prev value of immediate_quit.
Richard M. Stallman <rms@gnu.org>
parents: 20545
diff changeset
788 immediate_quit = old_immediate_quit;
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
789 start_byte = start_byte + cursor - base + 1;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
790 start = BYTE_TO_CHAR (start_byte);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
791 TEMP_SET_PT_BOTH (start, start_byte);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
792 return 0;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
793 }
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
794 else
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
795 if (++cursor == ceiling_addr)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
796 break;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
797 }
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
798 else
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
799 break;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
800 }
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
801 start_byte += cursor - base;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
802 }
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
803 }
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
804 else
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
805 {
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
806 while (start_byte > limit_byte)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
807 {
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
808 ceiling = BUFFER_FLOOR_OF (start_byte - 1);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
809 ceiling = max (limit_byte, ceiling);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
810 ceiling_addr = BYTE_POS_ADDR (ceiling) - 1;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
811 base = (cursor = BYTE_POS_ADDR (start_byte - 1) + 1);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
812 while (1)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
813 {
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
814 while (--cursor != ceiling_addr && *cursor != '\n')
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
815 ;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
816
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
817 if (cursor != ceiling_addr)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
818 {
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
819 if (++count == 0)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
820 {
20547
07053199a368 (scan_newline): Always restore prev value of immediate_quit.
Richard M. Stallman <rms@gnu.org>
parents: 20545
diff changeset
821 immediate_quit = old_immediate_quit;
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
822 /* Return the position AFTER the match we found. */
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
823 start_byte = start_byte + cursor - base + 1;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
824 start = BYTE_TO_CHAR (start_byte);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
825 TEMP_SET_PT_BOTH (start, start_byte);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
826 return 0;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
827 }
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
828 }
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
829 else
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
830 break;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
831 }
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
832 /* Here we add 1 to compensate for the last decrement
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
833 of CURSOR, which took it past the valid range. */
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
834 start_byte += cursor - base + 1;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
835 }
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
836 }
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
837
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
838 TEMP_SET_PT_BOTH (limit, limit_byte);
20547
07053199a368 (scan_newline): Always restore prev value of immediate_quit.
Richard M. Stallman <rms@gnu.org>
parents: 20545
diff changeset
839 immediate_quit = old_immediate_quit;
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
840
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
841 return count * direction;
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
842 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
843
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
844 int
7891
7d8e0f338e4a (find_next_newline_no_quit): New function.
Richard M. Stallman <rms@gnu.org>
parents: 7856
diff changeset
845 find_next_newline_no_quit (from, cnt)
7d8e0f338e4a (find_next_newline_no_quit): New function.
Richard M. Stallman <rms@gnu.org>
parents: 7856
diff changeset
846 register int from, cnt;
7d8e0f338e4a (find_next_newline_no_quit): New function.
Richard M. Stallman <rms@gnu.org>
parents: 7856
diff changeset
847 {
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
848 return scan_buffer ('\n', from, 0, cnt, (int *) 0, 0);
7891
7d8e0f338e4a (find_next_newline_no_quit): New function.
Richard M. Stallman <rms@gnu.org>
parents: 7856
diff changeset
849 }
7d8e0f338e4a (find_next_newline_no_quit): New function.
Richard M. Stallman <rms@gnu.org>
parents: 7856
diff changeset
850
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
851 /* Like find_next_newline, but returns position before the newline,
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
852 not after, and only search up to TO. This isn't just
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
853 find_next_newline (...)-1, because you might hit TO. */
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
854
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
855 int
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
856 find_before_next_newline (from, to, cnt)
9452
76f75b9091f1 (scan_buffer): After temporarily turning immediate_quit off, turn it
Jim Blandy <jimb@redhat.com>
parents: 9410
diff changeset
857 int from, to, cnt;
9410
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
858 {
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
859 int shortage;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
860 int pos = scan_buffer ('\n', from, to, cnt, &shortage, 1);
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
861
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
862 if (shortage == 0)
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
863 pos--;
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
864
8598c3d6f2f0 * search.c: #include "region-cache.h".
Jim Blandy <jimb@redhat.com>
parents: 9319
diff changeset
865 return pos;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
866 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
867
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
868 /* Subroutines of Lisp buffer search functions. */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
869
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
870 static Lisp_Object
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
871 search_command (string, bound, noerror, count, direction, RE, posix)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
872 Lisp_Object string, bound, noerror, count;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
873 int direction;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
874 int RE;
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
875 int posix;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
876 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
877 register int np;
20824
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
878 int lim, lim_byte;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
879 int n = direction;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
880
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
881 if (!NILP (count))
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
882 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
883 CHECK_NUMBER (count, 3);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
884 n *= XINT (count);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
885 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
886
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
887 CHECK_STRING (string, 0);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
888 if (NILP (bound))
20824
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
889 {
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
890 if (n > 0)
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
891 lim = ZV, lim_byte = ZV_BYTE;
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
892 else
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
893 lim = BEGV, lim_byte = BEGV_BYTE;
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
894 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
895 else
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
896 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
897 CHECK_NUMBER_COERCE_MARKER (bound, 1);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
898 lim = XINT (bound);
16039
855c8d8ba0f0 Change all references from point to PT.
Karl Heuer <kwzh@gnu.org>
parents: 15667
diff changeset
899 if (n > 0 ? lim < PT : lim > PT)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
900 error ("Invalid search bound (wrong side of point)");
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
901 if (lim > ZV)
20824
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
902 lim = ZV, lim_byte = ZV_BYTE;
20924
eda7e44ef9d9 (search_command): Check LIM in valid range
Karl Heuer <kwzh@gnu.org>
parents: 20898
diff changeset
903 else if (lim < BEGV)
20824
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
904 lim = BEGV, lim_byte = BEGV_BYTE;
20924
eda7e44ef9d9 (search_command): Check LIM in valid range
Karl Heuer <kwzh@gnu.org>
parents: 20898
diff changeset
905 else
eda7e44ef9d9 (search_command): Check LIM in valid range
Karl Heuer <kwzh@gnu.org>
parents: 20898
diff changeset
906 lim_byte = CHAR_TO_BYTE (lim);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
907 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
908
20824
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
909 np = search_buffer (string, PT, PT_BYTE, lim, lim_byte, n, RE,
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
910 (!NILP (current_buffer->case_fold_search)
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
911 ? current_buffer->case_canon_table
20875
4fac9830041a (search_command): Fix call to search_buffer.
Richard M. Stallman <rms@gnu.org>
parents: 20869
diff changeset
912 : Qnil),
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
913 (!NILP (current_buffer->case_fold_search)
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
914 ? current_buffer->case_eqv_table
20875
4fac9830041a (search_command): Fix call to search_buffer.
Richard M. Stallman <rms@gnu.org>
parents: 20869
diff changeset
915 : Qnil),
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
916 posix);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
917 if (np <= 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
918 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
919 if (NILP (noerror))
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
920 return signal_failure (string);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
921 if (!EQ (noerror, Qt))
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
922 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
923 if (lim < BEGV || lim > ZV)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
924 abort ();
20824
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
925 SET_PT_BOTH (lim, lim_byte);
1878
1c26d0049d4f (search_command): #if 0 previous change.
Richard M. Stallman <rms@gnu.org>
parents: 1877
diff changeset
926 return Qnil;
1c26d0049d4f (search_command): #if 0 previous change.
Richard M. Stallman <rms@gnu.org>
parents: 1877
diff changeset
927 #if 0 /* This would be clean, but maybe programs depend on
1c26d0049d4f (search_command): #if 0 previous change.
Richard M. Stallman <rms@gnu.org>
parents: 1877
diff changeset
928 a value of nil here. */
1877
7786f61ec635 (search_command): When moving to LIM on failure, return LIM.
Richard M. Stallman <rms@gnu.org>
parents: 1684
diff changeset
929 np = lim;
1878
1c26d0049d4f (search_command): #if 0 previous change.
Richard M. Stallman <rms@gnu.org>
parents: 1877
diff changeset
930 #endif
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
931 }
1877
7786f61ec635 (search_command): When moving to LIM on failure, return LIM.
Richard M. Stallman <rms@gnu.org>
parents: 1684
diff changeset
932 else
7786f61ec635 (search_command): When moving to LIM on failure, return LIM.
Richard M. Stallman <rms@gnu.org>
parents: 1684
diff changeset
933 return Qnil;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
934 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
935
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
936 if (np < BEGV || np > ZV)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
937 abort ();
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
938
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
939 SET_PT (np);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
940
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
941 return make_number (np);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
942 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
943
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
944 /* Return 1 if REGEXP it matches just one constant string. */
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
945
8950
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
946 static int
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
947 trivial_regexp_p (regexp)
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
948 Lisp_Object regexp;
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
949 {
21244
50929073a0ba Use STRING_BYTES and SET_STRING_BYTES.
Richard M. Stallman <rms@gnu.org>
parents: 21171
diff changeset
950 int len = STRING_BYTES (XSTRING (regexp));
8950
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
951 unsigned char *s = XSTRING (regexp)->data;
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
952 unsigned char c;
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
953 while (--len >= 0)
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
954 {
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
955 switch (*s++)
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
956 {
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
957 case '.': case '*': case '+': case '?': case '[': case '^': case '$':
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
958 return 0;
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
959 case '\\':
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
960 if (--len < 0)
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
961 return 0;
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
962 switch (*s++)
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
963 {
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
964 case '|': case '(': case ')': case '`': case '\'': case 'b':
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
965 case 'B': case '<': case '>': case 'w': case 'W': case 's':
12069
505dc29a68cf (trivial_regexp_p): = is special after \.
Karl Heuer <kwzh@gnu.org>
parents: 11678
diff changeset
966 case 'S': case '=':
17053
df904355f033 Include category.h and charset.h.
Karl Heuer <kwzh@gnu.org>
parents: 16880
diff changeset
967 case 'c': case 'C': /* for categoryspec and notcategoryspec */
12069
505dc29a68cf (trivial_regexp_p): = is special after \.
Karl Heuer <kwzh@gnu.org>
parents: 11678
diff changeset
968 case '1': case '2': case '3': case '4': case '5':
8950
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
969 case '6': case '7': case '8': case '9':
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
970 return 0;
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
971 }
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
972 }
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
973 }
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
974 return 1;
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
975 }
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
976
5556
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
977 /* Search for the n'th occurrence of STRING in the current buffer,
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
978 starting at position POS and stopping at position LIM,
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
979 treating STRING as a literal string if RE is false or as
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
980 a regular expression if RE is true.
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
981
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
982 If N is positive, searching is forward and LIM must be greater than POS.
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
983 If N is negative, searching is backward and LIM must be less than POS.
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
984
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
985 Returns -x if x occurrences remain to be found (x > 0),
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
986 or else the position at the beginning of the Nth occurrence
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
987 (if searching backward) or the end (if searching forward).
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
988
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
989 POSIX is nonzero if we want full backtracking (POSIX style)
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
990 for this pattern. 0 means backtrack only enough to get a valid match. */
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
991
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
992 #define TRANSLATE(out, trt, d) \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
993 do \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
994 { \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
995 if (! NILP (trt)) \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
996 { \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
997 Lisp_Object temp; \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
998 temp = Faref (trt, make_number (d)); \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
999 if (INTEGERP (temp)) \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1000 out = XINT (temp); \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1001 else \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1002 out = d; \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1003 } \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1004 else \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1005 out = d; \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1006 } \
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1007 while (0)
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1008
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
1009 static int
20824
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
1010 search_buffer (string, pos, pos_byte, lim, lim_byte, n,
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
1011 RE, trt, inverse_trt, posix)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1012 Lisp_Object string;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1013 int pos;
20824
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
1014 int pos_byte;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1015 int lim;
20824
97df0c6e753d (search_buffer): New args pos_byte and lim_byte.
Richard M. Stallman <rms@gnu.org>
parents: 20792
diff changeset
1016 int lim_byte;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1017 int n;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1018 int RE;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1019 Lisp_Object trt;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1020 Lisp_Object inverse_trt;
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
1021 int posix;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1022 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1023 int len = XSTRING (string)->size;
21244
50929073a0ba Use STRING_BYTES and SET_STRING_BYTES.
Richard M. Stallman <rms@gnu.org>
parents: 21171
diff changeset
1024 int len_byte = STRING_BYTES (XSTRING (string));
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1025 register int i;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1026
10032
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
1027 if (running_asynch_code)
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
1028 save_search_regs ();
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
1029
22082
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1030 /* Searching 0 times means don't move. */
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1031 /* Null string is found at starting position. */
22082
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1032 if (len == 0 || n == 0)
5556
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1033 {
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1034 set_search_regs (pos, 0);
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1035 return pos;
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1036 }
4299
7a2e1d7362c5 (search_buffer): If n is 0, just return POS.
Richard M. Stallman <rms@gnu.org>
parents: 3615
diff changeset
1037
8950
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1038 if (RE && !trivial_regexp_p (string))
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1039 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1040 unsigned char *p1, *p2;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1041 int s1, s2;
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
1042 struct re_pattern_buffer *bufp;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
1043
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1044 bufp = compile_pattern (string, &search_regs, trt, posix,
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1045 !NILP (current_buffer->enable_multibyte_characters));
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1046
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1047 immediate_quit = 1; /* Quit immediately if user types ^G,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1048 because letting this function finish
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1049 can take too long. */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1050 QUIT; /* Do a pending quit right away,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1051 to avoid paradoxical behavior */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1052 /* Get pointers and sizes of the two strings
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1053 that make up the visible portion of the buffer. */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1054
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1055 p1 = BEGV_ADDR;
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1056 s1 = GPT_BYTE - BEGV_BYTE;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1057 p2 = GAP_END_ADDR;
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1058 s2 = ZV_BYTE - GPT_BYTE;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1059 if (s1 < 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1060 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1061 p2 = p1;
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1062 s2 = ZV_BYTE - BEGV_BYTE;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1063 s1 = 0;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1064 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1065 if (s2 < 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1066 {
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1067 s1 = ZV_BYTE - BEGV_BYTE;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1068 s2 = 0;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1069 }
17463
bb9ae80d22e2 (looking_at_1): Set re_match_object.
Richard M. Stallman <rms@gnu.org>
parents: 17284
diff changeset
1070 re_match_object = Qnil;
bb9ae80d22e2 (looking_at_1): Set re_match_object.
Richard M. Stallman <rms@gnu.org>
parents: 17284
diff changeset
1071
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1072 while (n < 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1073 {
2475
052bbdf1b817 (search_buffer): Fix typo in previous change.
Richard M. Stallman <rms@gnu.org>
parents: 2439
diff changeset
1074 int val;
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
1075 val = re_search_2 (bufp, (char *) p1, s1, (char *) p2, s2,
20792
f0aa5cc14e8a (fast_string_match): Give re_search byte size of
Kenichi Handa <handa@m17n.org>
parents: 20706
diff changeset
1076 pos_byte - BEGV_BYTE, lim_byte - pos_byte,
f0aa5cc14e8a (fast_string_match): Give re_search byte size of
Kenichi Handa <handa@m17n.org>
parents: 20706
diff changeset
1077 &search_regs,
2475
052bbdf1b817 (search_buffer): Fix typo in previous change.
Richard M. Stallman <rms@gnu.org>
parents: 2439
diff changeset
1078 /* Don't allow match past current point */
20792
f0aa5cc14e8a (fast_string_match): Give re_search byte size of
Kenichi Handa <handa@m17n.org>
parents: 20706
diff changeset
1079 pos_byte - BEGV_BYTE);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1080 if (val == -2)
8950
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1081 {
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1082 matcher_overflow ();
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1083 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1084 if (val >= 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1085 {
20927
765fdbf766e4 (search_buffer): Update POS_BYTE for regexp search.
Kenichi Handa <handa@m17n.org>
parents: 20924
diff changeset
1086 pos_byte = search_regs.start[0] + BEGV_BYTE;
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
1087 for (i = 0; i < search_regs.num_regs; i++)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1088 if (search_regs.start[i] >= 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1089 {
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1090 search_regs.start[i]
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1091 = BYTE_TO_CHAR (search_regs.start[i] + BEGV_BYTE);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1092 search_regs.end[i]
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1093 = BYTE_TO_CHAR (search_regs.end[i] + BEGV_BYTE);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1094 }
9278
f2138d548313 (Flooking_at, skip_chars, search_buffer, set_search_regs, Fstore_match_data):
Karl Heuer <kwzh@gnu.org>
parents: 9113
diff changeset
1095 XSETBUFFER (last_thing_searched, current_buffer);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1096 /* Set pos to the new position. */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1097 pos = search_regs.start[0];
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1098 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1099 else
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1100 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1101 immediate_quit = 0;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1102 return (n);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1103 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1104 n++;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1105 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1106 while (n > 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1107 {
2475
052bbdf1b817 (search_buffer): Fix typo in previous change.
Richard M. Stallman <rms@gnu.org>
parents: 2439
diff changeset
1108 int val;
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
1109 val = re_search_2 (bufp, (char *) p1, s1, (char *) p2, s2,
20792
f0aa5cc14e8a (fast_string_match): Give re_search byte size of
Kenichi Handa <handa@m17n.org>
parents: 20706
diff changeset
1110 pos_byte - BEGV_BYTE, lim_byte - pos_byte,
f0aa5cc14e8a (fast_string_match): Give re_search byte size of
Kenichi Handa <handa@m17n.org>
parents: 20706
diff changeset
1111 &search_regs,
f0aa5cc14e8a (fast_string_match): Give re_search byte size of
Kenichi Handa <handa@m17n.org>
parents: 20706
diff changeset
1112 lim_byte - BEGV_BYTE);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1113 if (val == -2)
8950
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1114 {
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1115 matcher_overflow ();
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1116 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1117 if (val >= 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1118 {
20927
765fdbf766e4 (search_buffer): Update POS_BYTE for regexp search.
Kenichi Handa <handa@m17n.org>
parents: 20924
diff changeset
1119 pos_byte = search_regs.end[0] + BEGV_BYTE;
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
1120 for (i = 0; i < search_regs.num_regs; i++)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1121 if (search_regs.start[i] >= 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1122 {
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1123 search_regs.start[i]
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1124 = BYTE_TO_CHAR (search_regs.start[i] + BEGV_BYTE);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1125 search_regs.end[i]
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1126 = BYTE_TO_CHAR (search_regs.end[i] + BEGV_BYTE);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1127 }
9278
f2138d548313 (Flooking_at, skip_chars, search_buffer, set_search_regs, Fstore_match_data):
Karl Heuer <kwzh@gnu.org>
parents: 9113
diff changeset
1128 XSETBUFFER (last_thing_searched, current_buffer);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1129 pos = search_regs.end[0];
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1130 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1131 else
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1132 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1133 immediate_quit = 0;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1134 return (0 - n);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1135 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1136 n--;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1137 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1138 immediate_quit = 0;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1139 return (pos);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1140 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1141 else /* non-RE case */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1142 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1143 unsigned char *raw_pattern, *pat;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1144 int raw_pattern_size;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1145 int raw_pattern_size_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1146 unsigned char *patbuf;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1147 int multibyte = !NILP (current_buffer->enable_multibyte_characters);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1148 unsigned char *base_pat = XSTRING (string)->data;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1149 int charset_base = -1;
23876
8d2f38338c81 (search_buffer): Don't use Boyer-Moore
Kenichi Handa <handa@m17n.org>
parents: 23790
diff changeset
1150 int boyer_moore_ok = 1;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1151
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1152 /* MULTIBYTE says whether the text to be searched is multibyte.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1153 We must convert PATTERN to match that, or we will not really
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1154 find things right. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1155
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1156 if (multibyte == STRING_MULTIBYTE (string))
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1157 {
21915
8f1159b417c2 (search_buffer): Fix casts when assigning raw_pattern.
Richard M. Stallman <rms@gnu.org>
parents: 21887
diff changeset
1158 raw_pattern = (unsigned char *) XSTRING (string)->data;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1159 raw_pattern_size = XSTRING (string)->size;
21244
50929073a0ba Use STRING_BYTES and SET_STRING_BYTES.
Richard M. Stallman <rms@gnu.org>
parents: 21171
diff changeset
1160 raw_pattern_size_byte = STRING_BYTES (XSTRING (string));
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1161 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1162 else if (multibyte)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1163 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1164 raw_pattern_size = XSTRING (string)->size;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1165 raw_pattern_size_byte
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1166 = count_size_as_multibyte (XSTRING (string)->data,
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1167 raw_pattern_size);
21915
8f1159b417c2 (search_buffer): Fix casts when assigning raw_pattern.
Richard M. Stallman <rms@gnu.org>
parents: 21887
diff changeset
1168 raw_pattern = (unsigned char *) alloca (raw_pattern_size_byte + 1);
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1169 copy_text (XSTRING (string)->data, raw_pattern,
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1170 XSTRING (string)->size, 0, 1);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1171 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1172 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1173 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1174 /* Converting multibyte to single-byte.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1175
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1176 ??? Perhaps this conversion should be done in a special way
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1177 by subtracting nonascii-insert-offset from each non-ASCII char,
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1178 so that only the multibyte chars which really correspond to
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1179 the chosen single-byte character set can possibly match. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1180 raw_pattern_size = XSTRING (string)->size;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1181 raw_pattern_size_byte = XSTRING (string)->size;
21915
8f1159b417c2 (search_buffer): Fix casts when assigning raw_pattern.
Richard M. Stallman <rms@gnu.org>
parents: 21887
diff changeset
1182 raw_pattern = (unsigned char *) alloca (raw_pattern_size + 1);
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1183 copy_text (XSTRING (string)->data, raw_pattern,
21244
50929073a0ba Use STRING_BYTES and SET_STRING_BYTES.
Richard M. Stallman <rms@gnu.org>
parents: 21171
diff changeset
1184 STRING_BYTES (XSTRING (string)), 1, 0);
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1185 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1186
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1187 /* Copy and optionally translate the pattern. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1188 len = raw_pattern_size;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1189 len_byte = raw_pattern_size_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1190 patbuf = (unsigned char *) alloca (len_byte);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1191 pat = patbuf;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1192 base_pat = raw_pattern;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1193 if (multibyte)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1194 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1195 while (--len >= 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1196 {
26869
cb8fbc50812f (search_buffer): Adjusted for the change of CHAR_STRING.
Kenichi Handa <handa@m17n.org>
parents: 26088
diff changeset
1197 unsigned char str[MAX_MULTIBYTE_LENGTH];
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1198 int c, translated, inverse;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1199 int in_charlen, charlen;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1200
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1201 /* If we got here and the RE flag is set, it's because we're
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1202 dealing with a regexp known to be trivial, so the backslash
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1203 just quotes the next character. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1204 if (RE && *base_pat == '\\')
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1205 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1206 len--;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1207 len_byte--;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1208 base_pat++;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1209 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1210
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1211 c = STRING_CHAR_AND_LENGTH (base_pat, len_byte, in_charlen);
23876
8d2f38338c81 (search_buffer): Don't use Boyer-Moore
Kenichi Handa <handa@m17n.org>
parents: 23790
diff changeset
1212
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1213 /* Translate the character, if requested. */
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1214 TRANSLATE (translated, trt, c);
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1215 /* If translation changed the byte-length, go back
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1216 to the original character. */
26869
cb8fbc50812f (search_buffer): Adjusted for the change of CHAR_STRING.
Kenichi Handa <handa@m17n.org>
parents: 26088
diff changeset
1217 charlen = CHAR_STRING (translated, str);
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1218 if (in_charlen != charlen)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1219 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1220 translated = c;
26869
cb8fbc50812f (search_buffer): Adjusted for the change of CHAR_STRING.
Kenichi Handa <handa@m17n.org>
parents: 26088
diff changeset
1221 charlen = CHAR_STRING (c, str);
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1222 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1223
24014
0997bcfd8827 (search_buffer): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 23876
diff changeset
1224 /* If we are searching for something strange,
0997bcfd8827 (search_buffer): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 23876
diff changeset
1225 an invalid multibyte code, don't use boyer-moore. */
0997bcfd8827 (search_buffer): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 23876
diff changeset
1226 if (! ASCII_BYTE_P (translated)
0997bcfd8827 (search_buffer): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 23876
diff changeset
1227 && (charlen == 1 /* 8bit code */
0997bcfd8827 (search_buffer): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 23876
diff changeset
1228 || charlen != in_charlen /* invalid multibyte code */
0997bcfd8827 (search_buffer): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 23876
diff changeset
1229 ))
0997bcfd8827 (search_buffer): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 23876
diff changeset
1230 boyer_moore_ok = 0;
0997bcfd8827 (search_buffer): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 23876
diff changeset
1231
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1232 TRANSLATE (inverse, inverse_trt, c);
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1233
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1234 /* Did this char actually get translated?
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1235 Would any other char get translated into it? */
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1236 if (translated != c || inverse != c)
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1237 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1238 /* Keep track of which character set row
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1239 contains the characters that need translation. */
24014
0997bcfd8827 (search_buffer): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 23876
diff changeset
1240 int charset_base_code = c & ~CHAR_FIELD3_MASK;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1241 if (charset_base == -1)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1242 charset_base = charset_base_code;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1243 else if (charset_base != charset_base_code)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1244 /* If two different rows appear, needing translation,
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1245 then we cannot use boyer_moore search. */
23876
8d2f38338c81 (search_buffer): Don't use Boyer-Moore
Kenichi Handa <handa@m17n.org>
parents: 23790
diff changeset
1246 boyer_moore_ok = 0;
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1247 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1248
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1249 /* Store this character into the translated pattern. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1250 bcopy (str, pat, charlen);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1251 pat += charlen;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1252 base_pat += in_charlen;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1253 len_byte -= in_charlen;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1254 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1255 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1256 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1257 {
23876
8d2f38338c81 (search_buffer): Don't use Boyer-Moore
Kenichi Handa <handa@m17n.org>
parents: 23790
diff changeset
1258 /* Unibyte buffer. */
8d2f38338c81 (search_buffer): Don't use Boyer-Moore
Kenichi Handa <handa@m17n.org>
parents: 23790
diff changeset
1259 charset_base = 0;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1260 while (--len >= 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1261 {
23876
8d2f38338c81 (search_buffer): Don't use Boyer-Moore
Kenichi Handa <handa@m17n.org>
parents: 23790
diff changeset
1262 int c, translated;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1263
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1264 /* If we got here and the RE flag is set, it's because we're
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1265 dealing with a regexp known to be trivial, so the backslash
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1266 just quotes the next character. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1267 if (RE && *base_pat == '\\')
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1268 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1269 len--;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1270 base_pat++;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1271 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1272 c = *base_pat++;
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1273 TRANSLATE (translated, trt, c);
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1274 *pat++ = translated;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1275 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1276 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1277
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1278 len_byte = pat - patbuf;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1279 len = raw_pattern_size;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1280 pat = base_pat = patbuf;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1281
23876
8d2f38338c81 (search_buffer): Don't use Boyer-Moore
Kenichi Handa <handa@m17n.org>
parents: 23790
diff changeset
1282 if (boyer_moore_ok)
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1283 return boyer_moore (n, pat, len, len_byte, trt, inverse_trt,
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1284 pos, pos_byte, lim, lim_byte,
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1285 charset_base);
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1286 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1287 return simple_search (n, pat, len, len_byte, trt,
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1288 pos, pos_byte, lim, lim_byte);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1289 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1290 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1291
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1292 /* Do a simple string search N times for the string PAT,
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1293 whose length is LEN/LEN_BYTE,
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1294 from buffer position POS/POS_BYTE until LIM/LIM_BYTE.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1295 TRT is the translation table.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1296
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1297 Return the character position where the match is found.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1298 Otherwise, if M matches remained to be found, return -M.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1299
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1300 This kind of search works regardless of what is in PAT and
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1301 regardless of what is in TRT. It is used in cases where
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1302 boyer_moore cannot work. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1303
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1304 static int
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1305 simple_search (n, pat, len, len_byte, trt, pos, pos_byte, lim, lim_byte)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1306 int n;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1307 unsigned char *pat;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1308 int len, len_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1309 Lisp_Object trt;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1310 int pos, pos_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1311 int lim, lim_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1312 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1313 int multibyte = ! NILP (current_buffer->enable_multibyte_characters);
21117
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1314 int forward = n > 0;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1315
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1316 if (lim > pos && multibyte)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1317 while (n > 0)
8950
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1318 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1319 while (1)
20671
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
1320 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1321 /* Try matching at position POS. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1322 int this_pos = pos;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1323 int this_pos_byte = pos_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1324 int this_len = len;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1325 int this_len_byte = len_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1326 unsigned char *p = pat;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1327 if (pos + len > lim)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1328 goto stop;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1329
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1330 while (this_len > 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1331 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1332 int charlen, buf_charlen;
21117
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1333 int pat_ch, buf_ch;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1334
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1335 pat_ch = STRING_CHAR_AND_LENGTH (p, this_len_byte, charlen);
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1336 buf_ch = STRING_CHAR_AND_LENGTH (BYTE_POS_ADDR (this_pos_byte),
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1337 ZV_BYTE - this_pos_byte,
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1338 buf_charlen);
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1339 TRANSLATE (buf_ch, trt, buf_ch);
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1340
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1341 if (buf_ch != pat_ch)
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1342 break;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1343
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1344 this_len_byte -= charlen;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1345 this_len--;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1346 p += charlen;
20671
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
1347
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1348 this_pos_byte += buf_charlen;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1349 this_pos++;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1350 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1351
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1352 if (this_len == 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1353 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1354 pos += len;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1355 pos_byte += len_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1356 break;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1357 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1358
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1359 INC_BOTH (pos, pos_byte);
20671
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
1360 }
be91d6130341 (compile_pattern_1): If representation of STRING
Karl Heuer <kwzh@gnu.org>
parents: 20588
diff changeset
1361
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1362 n--;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1363 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1364 else if (lim > pos)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1365 while (n > 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1366 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1367 while (1)
8950
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1368 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1369 /* Try matching at position POS. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1370 int this_pos = pos;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1371 int this_len = len;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1372 unsigned char *p = pat;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1373
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1374 if (pos + len > lim)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1375 goto stop;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1376
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1377 while (this_len > 0)
8950
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1378 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1379 int pat_ch = *p++;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1380 int buf_ch = FETCH_BYTE (this_pos);
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1381 TRANSLATE (buf_ch, trt, buf_ch);
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1382
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1383 if (buf_ch != pat_ch)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1384 break;
21117
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1385
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1386 this_len--;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1387 this_pos++;
8950
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1388 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1389
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1390 if (this_len == 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1391 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1392 pos += len;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1393 break;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1394 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1395
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1396 pos++;
8950
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1397 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1398
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1399 n--;
8950
68ad2f08d735 (trivial_regexp_p): New function.
Karl Heuer <kwzh@gnu.org>
parents: 8584
diff changeset
1400 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1401 /* Backwards search. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1402 else if (lim < pos && multibyte)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1403 while (n < 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1404 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1405 while (1)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1406 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1407 /* Try matching at position POS. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1408 int this_pos = pos - len;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1409 int this_pos_byte = pos_byte - len_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1410 int this_len = len;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1411 int this_len_byte = len_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1412 unsigned char *p = pat;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1413
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1414 if (pos - len < lim)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1415 goto stop;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1416
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1417 while (this_len > 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1418 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1419 int charlen, buf_charlen;
21117
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1420 int pat_ch, buf_ch;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1421
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1422 pat_ch = STRING_CHAR_AND_LENGTH (p, this_len_byte, charlen);
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1423 buf_ch = STRING_CHAR_AND_LENGTH (BYTE_POS_ADDR (this_pos_byte),
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1424 ZV_BYTE - this_pos_byte,
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1425 buf_charlen);
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1426 TRANSLATE (buf_ch, trt, buf_ch);
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1427
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1428 if (buf_ch != pat_ch)
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1429 break;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1430
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1431 this_len_byte -= charlen;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1432 this_len--;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1433 p += charlen;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1434 this_pos_byte += buf_charlen;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1435 this_pos++;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1436 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1437
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1438 if (this_len == 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1439 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1440 pos -= len;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1441 pos_byte -= len_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1442 break;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1443 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1444
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1445 DEC_BOTH (pos, pos_byte);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1446 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1447
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1448 n++;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1449 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1450 else if (lim < pos)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1451 while (n < 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1452 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1453 while (1)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1454 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1455 /* Try matching at position POS. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1456 int this_pos = pos - len;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1457 int this_len = len;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1458 unsigned char *p = pat;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1459
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1460 if (pos - len < lim)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1461 goto stop;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1462
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1463 while (this_len > 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1464 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1465 int pat_ch = *p++;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1466 int buf_ch = FETCH_BYTE (this_pos);
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1467 TRANSLATE (buf_ch, trt, buf_ch);
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1468
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1469 if (buf_ch != pat_ch)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1470 break;
21117
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1471 this_len--;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1472 this_pos++;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1473 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1474
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1475 if (this_len == 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1476 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1477 pos -= len;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1478 break;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1479 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1480
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1481 pos--;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1482 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1483
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1484 n++;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1485 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1486
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1487 stop:
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1488 if (n == 0)
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1489 {
21117
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1490 if (forward)
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1491 set_search_regs ((multibyte ? pos_byte : pos) - len_byte, len_byte);
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1492 else
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1493 set_search_regs (multibyte ? pos_byte : pos, len_byte);
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1494
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1495 return pos;
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1496 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1497 else if (n > 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1498 return -n;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1499 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1500 return n;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1501 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1502
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1503 /* Do Boyer-Moore search N times for the string PAT,
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1504 whose length is LEN/LEN_BYTE,
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1505 from buffer position POS/POS_BYTE until LIM/LIM_BYTE.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1506 DIRECTION says which direction we search in.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1507 TRT and INVERSE_TRT are translation tables.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1508
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1509 This kind of search works if all the characters in PAT that have
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1510 nontrivial translation are the same aside from the last byte. This
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1511 makes it possible to translate just the last byte of a character,
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1512 and do so after just a simple test of the context.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1513
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1514 If that criterion is not satisfied, do not call this function. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1515
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1516 static int
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1517 boyer_moore (n, base_pat, len, len_byte, trt, inverse_trt,
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1518 pos, pos_byte, lim, lim_byte, charset_base)
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1519 int n;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1520 unsigned char *base_pat;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1521 int len, len_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1522 Lisp_Object trt;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1523 Lisp_Object inverse_trt;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1524 int pos, pos_byte;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1525 int lim, lim_byte;
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1526 int charset_base;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1527 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1528 int direction = ((n > 0) ? 1 : -1);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1529 register int dirlen;
31829
43566b0aec59 Avoid some more compiler warnings.
Gerd Moellmann <gerd@gnu.org>
parents: 31486
diff changeset
1530 int infinity, limit, k, stride_for_teases = 0;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1531 register int *BM_tab;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1532 int *BM_tab_base;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1533 register unsigned char *cursor, *p_limit;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1534 register int i, j;
21945
bda081af77e7 (boyer_moore): Check more reliably for ptr[1] being
Richard M. Stallman <rms@gnu.org>
parents: 21915
diff changeset
1535 unsigned char *pat, *pat_end;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1536 int multibyte = ! NILP (current_buffer->enable_multibyte_characters);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1537
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1538 unsigned char simple_translate[0400];
31829
43566b0aec59 Avoid some more compiler warnings.
Gerd Moellmann <gerd@gnu.org>
parents: 31486
diff changeset
1539 int translate_prev_byte = 0;
43566b0aec59 Avoid some more compiler warnings.
Gerd Moellmann <gerd@gnu.org>
parents: 31486
diff changeset
1540 int translate_anteprev_byte = 0;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1541
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1542 #ifdef C_ALLOCA
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1543 int BM_tab_space[0400];
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1544 BM_tab = &BM_tab_space[0];
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1545 #else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1546 BM_tab = (int *) alloca (0400 * sizeof (int));
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1547 #endif
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1548 /* The general approach is that we are going to maintain that we know */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1549 /* the first (closest to the present position, in whatever direction */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1550 /* we're searching) character that could possibly be the last */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1551 /* (furthest from present position) character of a valid match. We */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1552 /* advance the state of our knowledge by looking at that character */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1553 /* and seeing whether it indeed matches the last character of the */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1554 /* pattern. If it does, we take a closer look. If it does not, we */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1555 /* move our pointer (to putative last characters) as far as is */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1556 /* logically possible. This amount of movement, which I call a */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1557 /* stride, will be the length of the pattern if the actual character */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1558 /* appears nowhere in the pattern, otherwise it will be the distance */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1559 /* from the last occurrence of that character to the end of the */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1560 /* pattern. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1561 /* As a coding trick, an enormous stride is coded into the table for */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1562 /* characters that match the last character. This allows use of only */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1563 /* a single test, a test for having gone past the end of the */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1564 /* permissible match region, to test for both possible matches (when */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1565 /* the stride goes past the end immediately) and failure to */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1566 /* match (where you get nudged past the end one stride at a time). */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1567
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1568 /* Here we make a "mickey mouse" BM table. The stride of the search */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1569 /* is determined only by the last character of the putative match. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1570 /* If that character does not match, we will stride the proper */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1571 /* distance to propose a match that superimposes it on the last */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1572 /* instance of a character that matches it (per trt), or misses */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1573 /* it entirely if there is none. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1574
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1575 dirlen = len_byte * direction;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1576 infinity = dirlen - (lim_byte + pos_byte + len_byte + len_byte) * direction;
21945
bda081af77e7 (boyer_moore): Check more reliably for ptr[1] being
Richard M. Stallman <rms@gnu.org>
parents: 21915
diff changeset
1577
bda081af77e7 (boyer_moore): Check more reliably for ptr[1] being
Richard M. Stallman <rms@gnu.org>
parents: 21915
diff changeset
1578 /* Record position after the end of the pattern. */
bda081af77e7 (boyer_moore): Check more reliably for ptr[1] being
Richard M. Stallman <rms@gnu.org>
parents: 21915
diff changeset
1579 pat_end = base_pat + len_byte;
bda081af77e7 (boyer_moore): Check more reliably for ptr[1] being
Richard M. Stallman <rms@gnu.org>
parents: 21915
diff changeset
1580 /* BASE_PAT points to a character that we start scanning from.
bda081af77e7 (boyer_moore): Check more reliably for ptr[1] being
Richard M. Stallman <rms@gnu.org>
parents: 21915
diff changeset
1581 It is the first character in a forward search,
bda081af77e7 (boyer_moore): Check more reliably for ptr[1] being
Richard M. Stallman <rms@gnu.org>
parents: 21915
diff changeset
1582 the last character in a backward search. */
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1583 if (direction < 0)
21945
bda081af77e7 (boyer_moore): Check more reliably for ptr[1] being
Richard M. Stallman <rms@gnu.org>
parents: 21915
diff changeset
1584 base_pat = pat_end - 1;
bda081af77e7 (boyer_moore): Check more reliably for ptr[1] being
Richard M. Stallman <rms@gnu.org>
parents: 21915
diff changeset
1585
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1586 BM_tab_base = BM_tab;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1587 BM_tab += 0400;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1588 j = dirlen; /* to get it in a register */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1589 /* A character that does not appear in the pattern induces a */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1590 /* stride equal to the pattern length. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1591 while (BM_tab_base != BM_tab)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1592 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1593 *--BM_tab = j;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1594 *--BM_tab = j;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1595 *--BM_tab = j;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1596 *--BM_tab = j;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1597 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1598
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1599 /* We use this for translation, instead of TRT itself.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1600 We fill this in to handle the characters that actually
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1601 occur in the pattern. Others don't matter anyway! */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1602 bzero (simple_translate, sizeof simple_translate);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1603 for (i = 0; i < 0400; i++)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1604 simple_translate[i] = i;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1605
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1606 i = 0;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1607 while (i != infinity)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1608 {
21945
bda081af77e7 (boyer_moore): Check more reliably for ptr[1] being
Richard M. Stallman <rms@gnu.org>
parents: 21915
diff changeset
1609 unsigned char *ptr = base_pat + i;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1610 i += direction;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1611 if (i == dirlen)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1612 i = infinity;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1613 if (! NILP (trt))
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1614 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1615 int ch;
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1616 int untranslated;
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1617 int this_translated = 1;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1618
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1619 if (multibyte
21945
bda081af77e7 (boyer_moore): Check more reliably for ptr[1] being
Richard M. Stallman <rms@gnu.org>
parents: 21915
diff changeset
1620 /* Is *PTR the last byte of a character? */
bda081af77e7 (boyer_moore): Check more reliably for ptr[1] being
Richard M. Stallman <rms@gnu.org>
parents: 21915
diff changeset
1621 && (pat_end - ptr == 1 || CHAR_HEAD_P (ptr[1])))
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1622 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1623 unsigned char *charstart = ptr;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1624 while (! CHAR_HEAD_P (*charstart))
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1625 charstart--;
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1626 untranslated = STRING_CHAR (charstart, ptr - charstart + 1);
24716
ceeb1f1d1a88 (boyer_moore): Get charset base value of `untranslated'
Kenichi Handa <handa@m17n.org>
parents: 24433
diff changeset
1627 if (charset_base == (untranslated & ~CHAR_FIELD3_MASK))
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1628 {
21117
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1629 TRANSLATE (ch, trt, untranslated);
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1630 if (! CHAR_HEAD_P (*ptr))
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1631 {
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1632 translate_prev_byte = ptr[-1];
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1633 if (! CHAR_HEAD_P (translate_prev_byte))
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1634 translate_anteprev_byte = ptr[-2];
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1635 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1636 }
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1637 else
21117
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1638 {
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1639 this_translated = 0;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1640 ch = *ptr;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1641 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1642 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1643 else if (!multibyte)
20898
f69969e35e78 (simple_search): Call set_search_regs.
Richard M. Stallman <rms@gnu.org>
parents: 20875
diff changeset
1644 TRANSLATE (ch, trt, *ptr);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1645 else
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1646 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1647 ch = *ptr;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1648 this_translated = 0;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1649 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1650
21117
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1651 if (ch > 0400)
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1652 j = ((unsigned char) ch) | 0200;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1653 else
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1654 j = (unsigned char) ch;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1655
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1656 if (i == infinity)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1657 stride_for_teases = BM_tab[j];
21117
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1658
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1659 BM_tab[j] = dirlen - i;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1660 /* A translation table is accompanied by its inverse -- see */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1661 /* comment following downcase_table for details */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1662 if (this_translated)
21117
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1663 {
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1664 int starting_ch = ch;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1665 int starting_j = j;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1666 while (1)
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1667 {
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1668 TRANSLATE (ch, inverse_trt, ch);
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1669 if (ch > 0400)
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1670 j = ((unsigned char) ch) | 0200;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1671 else
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1672 j = (unsigned char) ch;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1673
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1674 /* For all the characters that map into CH,
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1675 set up simple_translate to map the last byte
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1676 into STARTING_J. */
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1677 simple_translate[j] = starting_j;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1678 if (ch == starting_ch)
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1679 break;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1680 BM_tab[j] = dirlen - i;
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1681 }
a88d2c555a06 (simple_search): Don't count a character until it matches!
Richard M. Stallman <rms@gnu.org>
parents: 20965
diff changeset
1682 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1683 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1684 else
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1685 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1686 j = *ptr;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1687
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1688 if (i == infinity)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1689 stride_for_teases = BM_tab[j];
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1690 BM_tab[j] = dirlen - i;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1691 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1692 /* stride_for_teases tells how much to stride if we get a */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1693 /* match on the far character but are subsequently */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1694 /* disappointed, by recording what the stride would have been */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1695 /* for that character if the last character had been */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1696 /* different. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1697 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1698 infinity = dirlen - infinity;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1699 pos_byte += dirlen - ((direction > 0) ? direction : 0);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1700 /* loop invariant - POS_BYTE points at where last char (first
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1701 char if reverse) of pattern would align in a possible match. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1702 while (n != 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1703 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1704 int tail_end;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1705 unsigned char *tail_end_ptr;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1706
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1707 /* It's been reported that some (broken) compiler thinks that
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1708 Boolean expressions in an arithmetic context are unsigned.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1709 Using an explicit ?1:0 prevents this. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1710 if ((lim_byte - pos_byte - ((direction > 0) ? 1 : 0)) * direction
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1711 < 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1712 return (n * (0 - direction));
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1713 /* First we do the part we can by pointers (maybe nothing) */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1714 QUIT;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1715 pat = base_pat;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1716 limit = pos_byte - dirlen + direction;
21457
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1717 if (direction > 0)
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1718 {
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1719 limit = BUFFER_CEILING_OF (limit);
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1720 /* LIMIT is now the last (not beyond-last!) value POS_BYTE
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1721 can take on without hitting edge of buffer or the gap. */
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1722 limit = min (limit, pos_byte + 20000);
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1723 limit = min (limit, lim_byte - 1);
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1724 }
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1725 else
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1726 {
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1727 limit = BUFFER_FLOOR_OF (limit);
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1728 /* LIMIT is now the last (not beyond-last!) value POS_BYTE
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1729 can take on without hitting edge of buffer or the gap. */
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1730 limit = max (limit, pos_byte - 20000);
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1731 limit = max (limit, lim_byte);
8c6ea32aadfa (min, max): Make these macros, not functions.
Karl Heuer <kwzh@gnu.org>
parents: 21248
diff changeset
1732 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1733 tail_end = BUFFER_CEILING_OF (pos_byte) + 1;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1734 tail_end_ptr = BYTE_POS_ADDR (tail_end);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1735
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1736 if ((limit - pos_byte) * direction > 20)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1737 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1738 unsigned char *p2;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1739
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1740 p_limit = BYTE_POS_ADDR (limit);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1741 p2 = (cursor = BYTE_POS_ADDR (pos_byte));
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1742 /* In this loop, pos + cursor - p2 is the surrogate for pos */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1743 while (1) /* use one cursor setting as long as i can */
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1744 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1745 if (direction > 0) /* worth duplicating */
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1746 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1747 /* Use signed comparison if appropriate
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1748 to make cursor+infinity sure to be > p_limit.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1749 Assuming that the buffer lies in a range of addresses
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1750 that are all "positive" (as ints) or all "negative",
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1751 either kind of comparison will work as long
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1752 as we don't step by infinity. So pick the kind
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1753 that works when we do step by infinity. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1754 if ((EMACS_INT) (p_limit + infinity) > (EMACS_INT) p_limit)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1755 while ((EMACS_INT) cursor <= (EMACS_INT) p_limit)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1756 cursor += BM_tab[*cursor];
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1757 else
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1758 while ((EMACS_UINT) cursor <= (EMACS_UINT) p_limit)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1759 cursor += BM_tab[*cursor];
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1760 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1761 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1762 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1763 if ((EMACS_INT) (p_limit + infinity) < (EMACS_INT) p_limit)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1764 while ((EMACS_INT) cursor >= (EMACS_INT) p_limit)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1765 cursor += BM_tab[*cursor];
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1766 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1767 while ((EMACS_UINT) cursor >= (EMACS_UINT) p_limit)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1768 cursor += BM_tab[*cursor];
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1769 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1770 /* If you are here, cursor is beyond the end of the searched region. */
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1771 /* This can happen if you match on the far character of the pattern, */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1772 /* because the "stride" of that character is infinity, a number able */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1773 /* to throw you well beyond the end of the search. It can also */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1774 /* happen if you fail to match within the permitted region and would */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1775 /* otherwise try a character beyond that region */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1776 if ((cursor - p_limit) * direction <= len_byte)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1777 break; /* a small overrun is genuine */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1778 cursor -= infinity; /* large overrun = hit */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1779 i = dirlen - direction;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1780 if (! NILP (trt))
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1781 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1782 while ((i -= direction) + direction != 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1783 {
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1784 int ch;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1785 cursor -= direction;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1786 /* Translate only the last byte of a character. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1787 if (! multibyte
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1788 || ((cursor == tail_end_ptr
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1789 || CHAR_HEAD_P (cursor[1]))
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1790 && (CHAR_HEAD_P (cursor[0])
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1791 || (translate_prev_byte == cursor[-1]
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1792 && (CHAR_HEAD_P (translate_prev_byte)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1793 || translate_anteprev_byte == cursor[-2])))))
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1794 ch = simple_translate[*cursor];
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1795 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1796 ch = *cursor;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1797 if (pat[i] != ch)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1798 break;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1799 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1800 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1801 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1802 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1803 while ((i -= direction) + direction != 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1804 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1805 cursor -= direction;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1806 if (pat[i] != *cursor)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1807 break;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1808 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1809 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1810 cursor += dirlen - i - direction; /* fix cursor */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1811 if (i + direction == 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1812 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1813 int position;
708
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
1814
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1815 cursor -= direction;
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1816
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1817 position = pos_byte + cursor - p2 + ((direction > 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1818 ? 1 - len_byte : 0);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1819 set_search_regs (position, len_byte);
708
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
1820
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1821 if ((n -= direction) != 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1822 cursor += dirlen; /* to resume search */
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1823 else
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1824 return ((direction > 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1825 ? search_regs.end[0] : search_regs.start[0]);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1826 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1827 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1828 cursor += stride_for_teases; /* <sigh> we lose - */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1829 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1830 pos_byte += cursor - p2;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1831 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1832 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1833 /* Now we'll pick up a clump that has to be done the hard */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1834 /* way because it covers a discontinuity */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1835 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1836 limit = ((direction > 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1837 ? BUFFER_CEILING_OF (pos_byte - dirlen + 1)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1838 : BUFFER_FLOOR_OF (pos_byte - dirlen - 1));
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1839 limit = ((direction > 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1840 ? min (limit + len_byte, lim_byte - 1)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1841 : max (limit - len_byte, lim_byte));
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1842 /* LIMIT is now the last value POS_BYTE can have
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1843 and still be valid for a possible match. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1844 while (1)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1845 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1846 /* This loop can be coded for space rather than */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1847 /* speed because it will usually run only once. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1848 /* (the reach is at most len + 21, and typically */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1849 /* does not exceed len) */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1850 while ((limit - pos_byte) * direction >= 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1851 pos_byte += BM_tab[FETCH_BYTE (pos_byte)];
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1852 /* now run the same tests to distinguish going off the */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1853 /* end, a match or a phony match. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1854 if ((pos_byte - limit) * direction <= len_byte)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1855 break; /* ran off the end */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1856 /* Found what might be a match.
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1857 Set POS_BYTE back to last (first if reverse) pos. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1858 pos_byte -= infinity;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1859 i = dirlen - direction;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1860 while ((i -= direction) + direction != 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1861 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1862 int ch;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1863 unsigned char *ptr;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1864 pos_byte -= direction;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1865 ptr = BYTE_POS_ADDR (pos_byte);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1866 /* Translate only the last byte of a character. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1867 if (! multibyte
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1868 || ((ptr == tail_end_ptr
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1869 || CHAR_HEAD_P (ptr[1]))
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1870 && (CHAR_HEAD_P (ptr[0])
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1871 || (translate_prev_byte == ptr[-1]
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1872 && (CHAR_HEAD_P (translate_prev_byte)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1873 || translate_anteprev_byte == ptr[-2])))))
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1874 ch = simple_translate[*ptr];
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1875 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1876 ch = *ptr;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1877 if (pat[i] != ch)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1878 break;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1879 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1880 /* Above loop has moved POS_BYTE part or all the way
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1881 back to the first pos (last pos if reverse).
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1882 Set it once again at the last (first if reverse) char. */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1883 pos_byte += dirlen - i- direction;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1884 if (i + direction == 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1885 {
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1886 int position;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1887 pos_byte -= direction;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1888
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1889 position = pos_byte + ((direction > 0) ? 1 - len_byte : 0);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1890
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1891 set_search_regs (position, len_byte);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1892
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1893 if ((n -= direction) != 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1894 pos_byte += dirlen; /* to resume search */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1895 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1896 return ((direction > 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1897 ? search_regs.end[0] : search_regs.start[0]);
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1898 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1899 else
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1900 pos_byte += stride_for_teases;
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1901 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1902 }
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1903 /* We have done one clump. Can we continue? */
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1904 if ((lim_byte - pos_byte) * direction < 0)
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1905 return ((0 - n) * direction);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1906 }
20869
c9f608f889b4 (boyer_moore, simple_search): New subroutines.
Richard M. Stallman <rms@gnu.org>
parents: 20824
diff changeset
1907 return BYTE_TO_CHAR (pos_byte);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1908 }
5556
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1909
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1910 /* Record beginning BEG_BYTE and end BEG_BYTE + NBYTES
22082
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1911 for the overall match just found in the current buffer.
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1912 Also clear out the match data for registers 1 and up. */
5556
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1913
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1914 static void
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1915 set_search_regs (beg_byte, nbytes)
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1916 int beg_byte, nbytes;
5556
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1917 {
22082
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1918 int i;
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1919
5556
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1920 /* Make sure we have registers in which to store
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1921 the match position. */
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1922 if (search_regs.num_regs == 0)
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1923 {
10250
422c3b96efda (set_search_regs): Really set search_regs.start and .end.
Richard M. Stallman <rms@gnu.org>
parents: 10141
diff changeset
1924 search_regs.start = (regoff_t *) xmalloc (2 * sizeof (regoff_t));
422c3b96efda (set_search_regs): Really set search_regs.start and .end.
Richard M. Stallman <rms@gnu.org>
parents: 10141
diff changeset
1925 search_regs.end = (regoff_t *) xmalloc (2 * sizeof (regoff_t));
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
1926 search_regs.num_regs = 2;
5556
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1927 }
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1928
22082
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1929 /* Clear out the other registers. */
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1930 for (i = 1; i < search_regs.num_regs; i++)
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1931 {
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1932 search_regs.start[i] = -1;
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1933 search_regs.end[i] = -1;
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1934 }
84bcdbc46d71 (search_buffer): Set search regs for all success with an empty string.
Richard M. Stallman <rms@gnu.org>
parents: 21988
diff changeset
1935
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1936 search_regs.start[0] = BYTE_TO_CHAR (beg_byte);
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
1937 search_regs.end[0] = BYTE_TO_CHAR (beg_byte + nbytes);
9278
f2138d548313 (Flooking_at, skip_chars, search_buffer, set_search_regs, Fstore_match_data):
Karl Heuer <kwzh@gnu.org>
parents: 9113
diff changeset
1938 XSETBUFFER (last_thing_searched, current_buffer);
5556
14161cfec24a (set_search_regs): New subroutine.
Richard M. Stallman <rms@gnu.org>
parents: 4954
diff changeset
1939 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1940
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1941 /* Given a string of words separated by word delimiters,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1942 compute a regexp that matches those exact words
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1943 separated by arbitrary punctuation. */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1944
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1945 static Lisp_Object
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1946 wordify (string)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1947 Lisp_Object string;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1948 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1949 register unsigned char *p, *o;
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1950 register int i, i_byte, len, punct_count = 0, word_count = 0;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1951 Lisp_Object val;
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1952 int prev_c = 0;
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1953 int adjust;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1954
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1955 CHECK_STRING (string, 0);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1956 p = XSTRING (string)->data;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1957 len = XSTRING (string)->size;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1958
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1959 for (i = 0, i_byte = 0; i < len; )
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1960 {
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1961 int c;
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1962
29018
2f43c508a9b5 (wordify): Use FETCH_STRING_CHAR_ADVANCE
Kenichi Handa <handa@m17n.org>
parents: 28886
diff changeset
1963 FETCH_STRING_CHAR_ADVANCE (c, string, i, i_byte);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1964
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1965 if (SYNTAX (c) != Sword)
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1966 {
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1967 punct_count++;
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1968 if (i > 0 && SYNTAX (prev_c) == Sword)
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1969 word_count++;
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1970 }
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1971
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1972 prev_c = c;
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1973 }
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1974
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1975 if (SYNTAX (prev_c) == Sword)
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1976 word_count++;
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1977 if (!word_count)
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1978 return build_string ("");
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1979
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
1980 adjust = - punct_count + 5 * (word_count - 1) + 4;
22640
929ad308aba6 (wordify): Fix i_byte even in unibyte case for copy loop.
Richard M. Stallman <rms@gnu.org>
parents: 22533
diff changeset
1981 if (STRING_MULTIBYTE (string))
929ad308aba6 (wordify): Fix i_byte even in unibyte case for copy loop.
Richard M. Stallman <rms@gnu.org>
parents: 22533
diff changeset
1982 val = make_uninit_multibyte_string (len + adjust,
929ad308aba6 (wordify): Fix i_byte even in unibyte case for copy loop.
Richard M. Stallman <rms@gnu.org>
parents: 22533
diff changeset
1983 STRING_BYTES (XSTRING (string))
929ad308aba6 (wordify): Fix i_byte even in unibyte case for copy loop.
Richard M. Stallman <rms@gnu.org>
parents: 22533
diff changeset
1984 + adjust);
929ad308aba6 (wordify): Fix i_byte even in unibyte case for copy loop.
Richard M. Stallman <rms@gnu.org>
parents: 22533
diff changeset
1985 else
929ad308aba6 (wordify): Fix i_byte even in unibyte case for copy loop.
Richard M. Stallman <rms@gnu.org>
parents: 22533
diff changeset
1986 val = make_uninit_string (len + adjust);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1987
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1988 o = XSTRING (val)->data;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1989 *o++ = '\\';
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1990 *o++ = 'b';
21887
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
1991 prev_c = 0;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
1992
21887
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
1993 for (i = 0, i_byte = 0; i < len; )
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
1994 {
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
1995 int c;
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
1996 int i_byte_orig = i_byte;
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
1997
29018
2f43c508a9b5 (wordify): Use FETCH_STRING_CHAR_ADVANCE
Kenichi Handa <handa@m17n.org>
parents: 28886
diff changeset
1998 FETCH_STRING_CHAR_ADVANCE (c, string, i, i_byte);
21887
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
1999
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2000 if (SYNTAX (c) == Sword)
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2001 {
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2002 bcopy (&XSTRING (string)->data[i_byte_orig], o,
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2003 i_byte - i_byte_orig);
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2004 o += i_byte - i_byte_orig;
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2005 }
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2006 else if (i > 0 && SYNTAX (prev_c) == Sword && --word_count)
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2007 {
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2008 *o++ = '\\';
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2009 *o++ = 'W';
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2010 *o++ = '\\';
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2011 *o++ = 'W';
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2012 *o++ = '*';
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2013 }
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2014
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2015 prev_c = c;
1c9f20274f76 (wordify): Do the second loop by chars, not by bytes.
Richard M. Stallman <rms@gnu.org>
parents: 21531
diff changeset
2016 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2017
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2018 *o++ = '\\';
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2019 *o++ = 'b';
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2020
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2021 return val;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2022 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2023
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2024 DEFUN ("search-backward", Fsearch_backward, Ssearch_backward, 1, 4,
19541
e7876a076881 (Fsearch_backward): Inherit the current input method on
Kenichi Handa <handa@m17n.org>
parents: 18762
diff changeset
2025 "MSearch backward: ",
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2026 "Search backward from point for STRING.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2027 Set point to the beginning of the occurrence found, and return point.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2028 An optional second argument bounds the search; it is a buffer position.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2029 The match found must not extend before that position.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2030 Optional third argument, if t, means if fail just return nil (no error).\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2031 If not nil and not t, position at limit of search and return nil.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2032 Optional fourth argument is repeat count--search for successive occurrences.\n\
32387
7e7aced48811 (Fsearch_backward, Fsearch_forward): Doc fix.
Eli Zaretskii <eliz@gnu.org>
parents: 31829
diff changeset
2033 \n\
7e7aced48811 (Fsearch_backward, Fsearch_forward): Doc fix.
Eli Zaretskii <eliz@gnu.org>
parents: 31829
diff changeset
2034 Search case-sensitivity is determined by the value of the variable\n\
7e7aced48811 (Fsearch_backward, Fsearch_forward): Doc fix.
Eli Zaretskii <eliz@gnu.org>
parents: 31829
diff changeset
2035 `case-fold-search', which see.\n\
7e7aced48811 (Fsearch_backward, Fsearch_forward): Doc fix.
Eli Zaretskii <eliz@gnu.org>
parents: 31829
diff changeset
2036 \n\
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2037 See also the functions `match-beginning', `match-end' and `replace-match'.")
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2038 (string, bound, noerror, count)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2039 Lisp_Object string, bound, noerror, count;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2040 {
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2041 return search_command (string, bound, noerror, count, -1, 0, 0);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2042 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2043
19541
e7876a076881 (Fsearch_backward): Inherit the current input method on
Kenichi Handa <handa@m17n.org>
parents: 18762
diff changeset
2044 DEFUN ("search-forward", Fsearch_forward, Ssearch_forward, 1, 4, "MSearch: ",
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2045 "Search forward from point for STRING.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2046 Set point to the end of the occurrence found, and return point.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2047 An optional second argument bounds the search; it is a buffer position.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2048 The match found must not extend after that position. nil is equivalent\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2049 to (point-max).\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2050 Optional third argument, if t, means if fail just return nil (no error).\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2051 If not nil and not t, move to limit of search and return nil.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2052 Optional fourth argument is repeat count--search for successive occurrences.\n\
32387
7e7aced48811 (Fsearch_backward, Fsearch_forward): Doc fix.
Eli Zaretskii <eliz@gnu.org>
parents: 31829
diff changeset
2053 \n\
7e7aced48811 (Fsearch_backward, Fsearch_forward): Doc fix.
Eli Zaretskii <eliz@gnu.org>
parents: 31829
diff changeset
2054 Search case-sensitivity is determined by the value of the variable\n\
7e7aced48811 (Fsearch_backward, Fsearch_forward): Doc fix.
Eli Zaretskii <eliz@gnu.org>
parents: 31829
diff changeset
2055 `case-fold-search', which see.\n\
7e7aced48811 (Fsearch_backward, Fsearch_forward): Doc fix.
Eli Zaretskii <eliz@gnu.org>
parents: 31829
diff changeset
2056 \n\
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2057 See also the functions `match-beginning', `match-end' and `replace-match'.")
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2058 (string, bound, noerror, count)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2059 Lisp_Object string, bound, noerror, count;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2060 {
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2061 return search_command (string, bound, noerror, count, 1, 0, 0);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2062 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2063
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2064 DEFUN ("word-search-backward", Fword_search_backward, Sword_search_backward, 1, 4,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2065 "sWord search backward: ",
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2066 "Search backward from point for STRING, ignoring differences in punctuation.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2067 Set point to the beginning of the occurrence found, and return point.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2068 An optional second argument bounds the search; it is a buffer position.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2069 The match found must not extend before that position.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2070 Optional third argument, if t, means if fail just return nil (no error).\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2071 If not nil and not t, move to limit of search and return nil.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2072 Optional fourth argument is repeat count--search for successive occurrences.")
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2073 (string, bound, noerror, count)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2074 Lisp_Object string, bound, noerror, count;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2075 {
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2076 return search_command (wordify (string), bound, noerror, count, -1, 1, 0);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2077 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2078
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2079 DEFUN ("word-search-forward", Fword_search_forward, Sword_search_forward, 1, 4,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2080 "sWord search: ",
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2081 "Search forward from point for STRING, ignoring differences in punctuation.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2082 Set point to the end of the occurrence found, and return point.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2083 An optional second argument bounds the search; it is a buffer position.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2084 The match found must not extend after that position.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2085 Optional third argument, if t, means if fail just return nil (no error).\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2086 If not nil and not t, move to limit of search and return nil.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2087 Optional fourth argument is repeat count--search for successive occurrences.")
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2088 (string, bound, noerror, count)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2089 Lisp_Object string, bound, noerror, count;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2090 {
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2091 return search_command (wordify (string), bound, noerror, count, 1, 1, 0);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2092 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2093
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2094 DEFUN ("re-search-backward", Fre_search_backward, Sre_search_backward, 1, 4,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2095 "sRE search backward: ",
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2096 "Search backward from point for match for regular expression REGEXP.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2097 Set point to the beginning of the match, and return point.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2098 The match found is the one starting last in the buffer\n\
6297
b44907fd0ff0 (Fre_search_forward, Fre_search_backward): Doc fix.
Karl Heuer <kwzh@gnu.org>
parents: 6196
diff changeset
2099 and yet ending before the origin of the search.\n\
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2100 An optional second argument bounds the search; it is a buffer position.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2101 The match found must start at or after that position.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2102 Optional third argument, if t, means if fail just return nil (no error).\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2103 If not nil and not t, move to limit of search and return nil.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2104 Optional fourth argument is repeat count--search for successive occurrences.\n\
29335
1fa96662f8ca (Fre_search_forward, Fre_search_backward)
Jason Rumney <jasonr@gnu.org>
parents: 29300
diff changeset
2105 See also the functions `match-beginning', `match-end', `match-string',\n\
29300
505deeadeb1f (Fre_search_forward, Fre_search_backward)
Gerd Moellmann <gerd@gnu.org>
parents: 29018
diff changeset
2106 and `replace-match'.")
6297
b44907fd0ff0 (Fre_search_forward, Fre_search_backward): Doc fix.
Karl Heuer <kwzh@gnu.org>
parents: 6196
diff changeset
2107 (regexp, bound, noerror, count)
b44907fd0ff0 (Fre_search_forward, Fre_search_backward): Doc fix.
Karl Heuer <kwzh@gnu.org>
parents: 6196
diff changeset
2108 Lisp_Object regexp, bound, noerror, count;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2109 {
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2110 return search_command (regexp, bound, noerror, count, -1, 1, 0);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2111 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2112
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2113 DEFUN ("re-search-forward", Fre_search_forward, Sre_search_forward, 1, 4,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2114 "sRE search: ",
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2115 "Search forward from point for regular expression REGEXP.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2116 Set point to the end of the occurrence found, and return point.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2117 An optional second argument bounds the search; it is a buffer position.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2118 The match found must not extend after that position.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2119 Optional third argument, if t, means if fail just return nil (no error).\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2120 If not nil and not t, move to limit of search and return nil.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2121 Optional fourth argument is repeat count--search for successive occurrences.\n\
29335
1fa96662f8ca (Fre_search_forward, Fre_search_backward)
Jason Rumney <jasonr@gnu.org>
parents: 29300
diff changeset
2122 See also the functions `match-beginning', `match-end', `match-string',\n\
29300
505deeadeb1f (Fre_search_forward, Fre_search_backward)
Gerd Moellmann <gerd@gnu.org>
parents: 29018
diff changeset
2123 and `replace-match'.")
6297
b44907fd0ff0 (Fre_search_forward, Fre_search_backward): Doc fix.
Karl Heuer <kwzh@gnu.org>
parents: 6196
diff changeset
2124 (regexp, bound, noerror, count)
b44907fd0ff0 (Fre_search_forward, Fre_search_backward): Doc fix.
Karl Heuer <kwzh@gnu.org>
parents: 6196
diff changeset
2125 Lisp_Object regexp, bound, noerror, count;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2126 {
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2127 return search_command (regexp, bound, noerror, count, 1, 1, 0);
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2128 }
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2129
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2130 DEFUN ("posix-search-backward", Fposix_search_backward, Sposix_search_backward, 1, 4,
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2131 "sPosix search backward: ",
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2132 "Search backward from point for match for regular expression REGEXP.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2133 Find the longest match in accord with Posix regular expression rules.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2134 Set point to the beginning of the match, and return point.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2135 The match found is the one starting last in the buffer\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2136 and yet ending before the origin of the search.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2137 An optional second argument bounds the search; it is a buffer position.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2138 The match found must start at or after that position.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2139 Optional third argument, if t, means if fail just return nil (no error).\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2140 If not nil and not t, move to limit of search and return nil.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2141 Optional fourth argument is repeat count--search for successive occurrences.\n\
29335
1fa96662f8ca (Fre_search_forward, Fre_search_backward)
Jason Rumney <jasonr@gnu.org>
parents: 29300
diff changeset
2142 See also the functions `match-beginning', `match-end', `match-string',\n\
29300
505deeadeb1f (Fre_search_forward, Fre_search_backward)
Gerd Moellmann <gerd@gnu.org>
parents: 29018
diff changeset
2143 and `replace-match'.")
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2144 (regexp, bound, noerror, count)
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2145 Lisp_Object regexp, bound, noerror, count;
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2146 {
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2147 return search_command (regexp, bound, noerror, count, -1, 1, 1);
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2148 }
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2149
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2150 DEFUN ("posix-search-forward", Fposix_search_forward, Sposix_search_forward, 1, 4,
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2151 "sPosix search: ",
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2152 "Search forward from point for regular expression REGEXP.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2153 Find the longest match in accord with Posix regular expression rules.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2154 Set point to the end of the occurrence found, and return point.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2155 An optional second argument bounds the search; it is a buffer position.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2156 The match found must not extend after that position.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2157 Optional third argument, if t, means if fail just return nil (no error).\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2158 If not nil and not t, move to limit of search and return nil.\n\
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2159 Optional fourth argument is repeat count--search for successive occurrences.\n\
29335
1fa96662f8ca (Fre_search_forward, Fre_search_backward)
Jason Rumney <jasonr@gnu.org>
parents: 29300
diff changeset
2160 See also the functions `match-beginning', `match-end', `match-string',\n\
29300
505deeadeb1f (Fre_search_forward, Fre_search_backward)
Gerd Moellmann <gerd@gnu.org>
parents: 29018
diff changeset
2161 and `replace-match'.")
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2162 (regexp, bound, noerror, count)
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2163 Lisp_Object regexp, bound, noerror, count;
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2164 {
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2165 return search_command (regexp, bound, noerror, count, 1, 1, 1);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2166 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2167
12807
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2168 DEFUN ("replace-match", Freplace_match, Sreplace_match, 1, 5, 0,
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2169 "Replace text matched by last search with NEWTEXT.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2170 If second arg FIXEDCASE is non-nil, do not alter case of replacement text.\n\
6543
33032ee16c7c (Freplace_match): Doc fix.
Richard M. Stallman <rms@gnu.org>
parents: 6343
diff changeset
2171 Otherwise maybe capitalize the whole text, or maybe just word initials,\n\
33032ee16c7c (Freplace_match): Doc fix.
Richard M. Stallman <rms@gnu.org>
parents: 6343
diff changeset
2172 based on the replaced text.\n\
33032ee16c7c (Freplace_match): Doc fix.
Richard M. Stallman <rms@gnu.org>
parents: 6343
diff changeset
2173 If the replaced text has only capital letters\n\
33032ee16c7c (Freplace_match): Doc fix.
Richard M. Stallman <rms@gnu.org>
parents: 6343
diff changeset
2174 and has at least one multiletter word, convert NEWTEXT to all caps.\n\
33032ee16c7c (Freplace_match): Doc fix.
Richard M. Stallman <rms@gnu.org>
parents: 6343
diff changeset
2175 If the replaced text has at least one word starting with a capital letter,\n\
33032ee16c7c (Freplace_match): Doc fix.
Richard M. Stallman <rms@gnu.org>
parents: 6343
diff changeset
2176 then capitalize each word in NEWTEXT.\n\n\
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2177 If third arg LITERAL is non-nil, insert NEWTEXT literally.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2178 Otherwise treat `\\' as special:\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2179 `\\&' in NEWTEXT means substitute original matched text.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2180 `\\N' means substitute what matched the Nth `\\(...\\)'.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2181 If Nth parens didn't match, substitute nothing.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2182 `\\\\' means insert one `\\'.\n\
708
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2183 FIXEDCASE and LITERAL are optional arguments.\n\
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2184 Leaves point at end of replacement text.\n\
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2185 \n\
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2186 The optional fourth argument STRING can be a string to modify.\n\
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2187 In that case, this function creates and returns a new string\n\
12807
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2188 which is made by replacing the part of STRING that was matched.\n\
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2189 \n\
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2190 The optional fifth argument SUBEXP specifies a subexpression of the match.\n\
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2191 It says to replace just that subexpression instead of the whole match.\n\
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2192 This is useful only after a regular expression search or match\n\
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2193 since only regular expressions have distinguished subexpressions.")
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2194 (newtext, fixedcase, literal, string, subexp)
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2195 Lisp_Object newtext, fixedcase, literal, string, subexp;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2196 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2197 enum { nochange, all_caps, cap_initial } case_action;
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2198 register int pos, pos_byte;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2199 int some_multiletter_word;
2393
a35d2c5cbb3b (Freplace_match): Clean up criterion about converting case.
Richard M. Stallman <rms@gnu.org>
parents: 1926
diff changeset
2200 int some_lowercase;
7674
947d24fefd9e (Freplace_match): Improve capitalization heuristics.
Karl Heuer <kwzh@gnu.org>
parents: 7673
diff changeset
2201 int some_uppercase;
8526
2b7b23059f1b (Freplace_match): Treat caseless initial like a lowercase initial.
Richard M. Stallman <rms@gnu.org>
parents: 7891
diff changeset
2202 int some_nonuppercase_initial;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2203 register int c, prevc;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2204 int inslen;
12807
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2205 int sub;
18081
300068b4fcef (Freplace_match): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 18077
diff changeset
2206 int opoint, newpoint;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2207
4882
8c09f87f5087 (Freplace_match): Fix argument names to match doc string.
Brian Fox <bfox@gnu.org>
parents: 4832
diff changeset
2208 CHECK_STRING (newtext, 0);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2209
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2210 if (! NILP (string))
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2211 CHECK_STRING (string, 4);
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2212
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2213 case_action = nochange; /* We tried an initialization */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2214 /* but some C compilers blew it */
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2215
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2216 if (search_regs.num_regs <= 0)
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2217 error ("replace-match called before any match found");
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2218
12807
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2219 if (NILP (subexp))
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2220 sub = 0;
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2221 else
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2222 {
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2223 CHECK_NUMBER (subexp, 3);
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2224 sub = XINT (subexp);
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2225 if (sub < 0 || sub >= search_regs.num_regs)
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2226 args_out_of_range (subexp, make_number (search_regs.num_regs));
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2227 }
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2228
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2229 if (NILP (string))
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2230 {
12807
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2231 if (search_regs.start[sub] < BEGV
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2232 || search_regs.start[sub] > search_regs.end[sub]
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2233 || search_regs.end[sub] > ZV)
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2234 args_out_of_range (make_number (search_regs.start[sub]),
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2235 make_number (search_regs.end[sub]));
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2236 }
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2237 else
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2238 {
12807
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2239 if (search_regs.start[sub] < 0
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2240 || search_regs.start[sub] > search_regs.end[sub]
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2241 || search_regs.end[sub] > XSTRING (string)->size)
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2242 args_out_of_range (make_number (search_regs.start[sub]),
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2243 make_number (search_regs.end[sub]));
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2244 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2245
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2246 if (NILP (fixedcase))
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2247 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2248 /* Decide how to casify by examining the matched text. */
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2249 int last;
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2250
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2251 pos = search_regs.start[sub];
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2252 last = search_regs.end[sub];
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2253
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
2254 if (NILP (string))
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2255 pos_byte = CHAR_TO_BYTE (pos);
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
2256 else
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2257 pos_byte = string_char_to_byte (string, pos);
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
2258
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2259 prevc = '\n';
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2260 case_action = all_caps;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2261
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2262 /* some_multiletter_word is set nonzero if any original word
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2263 is more than one letter long. */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2264 some_multiletter_word = 0;
2393
a35d2c5cbb3b (Freplace_match): Clean up criterion about converting case.
Richard M. Stallman <rms@gnu.org>
parents: 1926
diff changeset
2265 some_lowercase = 0;
8526
2b7b23059f1b (Freplace_match): Treat caseless initial like a lowercase initial.
Richard M. Stallman <rms@gnu.org>
parents: 7891
diff changeset
2266 some_nonuppercase_initial = 0;
7674
947d24fefd9e (Freplace_match): Improve capitalization heuristics.
Karl Heuer <kwzh@gnu.org>
parents: 7673
diff changeset
2267 some_uppercase = 0;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2268
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2269 while (pos < last)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2270 {
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2271 if (NILP (string))
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2272 {
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2273 c = FETCH_CHAR (pos_byte);
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2274 INC_BOTH (pos, pos_byte);
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2275 }
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2276 else
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2277 FETCH_STRING_CHAR_ADVANCE (c, string, pos, pos_byte);
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2278
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2279 if (LOWERCASEP (c))
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2280 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2281 /* Cannot be all caps if any original char is lower case */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2282
2393
a35d2c5cbb3b (Freplace_match): Clean up criterion about converting case.
Richard M. Stallman <rms@gnu.org>
parents: 1926
diff changeset
2283 some_lowercase = 1;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2284 if (SYNTAX (prevc) != Sword)
8526
2b7b23059f1b (Freplace_match): Treat caseless initial like a lowercase initial.
Richard M. Stallman <rms@gnu.org>
parents: 7891
diff changeset
2285 some_nonuppercase_initial = 1;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2286 else
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2287 some_multiletter_word = 1;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2288 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2289 else if (!NOCASEP (c))
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2290 {
7674
947d24fefd9e (Freplace_match): Improve capitalization heuristics.
Karl Heuer <kwzh@gnu.org>
parents: 7673
diff changeset
2291 some_uppercase = 1;
2393
a35d2c5cbb3b (Freplace_match): Clean up criterion about converting case.
Richard M. Stallman <rms@gnu.org>
parents: 1926
diff changeset
2292 if (SYNTAX (prevc) != Sword)
6679
490b7e2db978 (Freplace_match): Don't capitalize unless all matched words are capitalized.
Karl Heuer <kwzh@gnu.org>
parents: 6543
diff changeset
2293 ;
2393
a35d2c5cbb3b (Freplace_match): Clean up criterion about converting case.
Richard M. Stallman <rms@gnu.org>
parents: 1926
diff changeset
2294 else
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2295 some_multiletter_word = 1;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2296 }
8526
2b7b23059f1b (Freplace_match): Treat caseless initial like a lowercase initial.
Richard M. Stallman <rms@gnu.org>
parents: 7891
diff changeset
2297 else
2b7b23059f1b (Freplace_match): Treat caseless initial like a lowercase initial.
Richard M. Stallman <rms@gnu.org>
parents: 7891
diff changeset
2298 {
2b7b23059f1b (Freplace_match): Treat caseless initial like a lowercase initial.
Richard M. Stallman <rms@gnu.org>
parents: 7891
diff changeset
2299 /* If the initial is a caseless word constituent,
2b7b23059f1b (Freplace_match): Treat caseless initial like a lowercase initial.
Richard M. Stallman <rms@gnu.org>
parents: 7891
diff changeset
2300 treat that like a lowercase initial. */
2b7b23059f1b (Freplace_match): Treat caseless initial like a lowercase initial.
Richard M. Stallman <rms@gnu.org>
parents: 7891
diff changeset
2301 if (SYNTAX (prevc) != Sword)
2b7b23059f1b (Freplace_match): Treat caseless initial like a lowercase initial.
Richard M. Stallman <rms@gnu.org>
parents: 7891
diff changeset
2302 some_nonuppercase_initial = 1;
2b7b23059f1b (Freplace_match): Treat caseless initial like a lowercase initial.
Richard M. Stallman <rms@gnu.org>
parents: 7891
diff changeset
2303 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2304
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2305 prevc = c;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2306 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2307
2393
a35d2c5cbb3b (Freplace_match): Clean up criterion about converting case.
Richard M. Stallman <rms@gnu.org>
parents: 1926
diff changeset
2308 /* Convert to all caps if the old text is all caps
a35d2c5cbb3b (Freplace_match): Clean up criterion about converting case.
Richard M. Stallman <rms@gnu.org>
parents: 1926
diff changeset
2309 and has at least one multiletter word. */
a35d2c5cbb3b (Freplace_match): Clean up criterion about converting case.
Richard M. Stallman <rms@gnu.org>
parents: 1926
diff changeset
2310 if (! some_lowercase && some_multiletter_word)
a35d2c5cbb3b (Freplace_match): Clean up criterion about converting case.
Richard M. Stallman <rms@gnu.org>
parents: 1926
diff changeset
2311 case_action = all_caps;
6679
490b7e2db978 (Freplace_match): Don't capitalize unless all matched words are capitalized.
Karl Heuer <kwzh@gnu.org>
parents: 6543
diff changeset
2312 /* Capitalize each word, if the old text has all capitalized words. */
8526
2b7b23059f1b (Freplace_match): Treat caseless initial like a lowercase initial.
Richard M. Stallman <rms@gnu.org>
parents: 7891
diff changeset
2313 else if (!some_nonuppercase_initial && some_multiletter_word)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2314 case_action = cap_initial;
8526
2b7b23059f1b (Freplace_match): Treat caseless initial like a lowercase initial.
Richard M. Stallman <rms@gnu.org>
parents: 7891
diff changeset
2315 else if (!some_nonuppercase_initial && some_uppercase)
7674
947d24fefd9e (Freplace_match): Improve capitalization heuristics.
Karl Heuer <kwzh@gnu.org>
parents: 7673
diff changeset
2316 /* Should x -> yz, operating on X, give Yz or YZ?
947d24fefd9e (Freplace_match): Improve capitalization heuristics.
Karl Heuer <kwzh@gnu.org>
parents: 7673
diff changeset
2317 We'll assume the latter. */
947d24fefd9e (Freplace_match): Improve capitalization heuristics.
Karl Heuer <kwzh@gnu.org>
parents: 7673
diff changeset
2318 case_action = all_caps;
2393
a35d2c5cbb3b (Freplace_match): Clean up criterion about converting case.
Richard M. Stallman <rms@gnu.org>
parents: 1926
diff changeset
2319 else
a35d2c5cbb3b (Freplace_match): Clean up criterion about converting case.
Richard M. Stallman <rms@gnu.org>
parents: 1926
diff changeset
2320 case_action = nochange;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2321 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2322
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2323 /* Do replacement in a string. */
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2324 if (!NILP (string))
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2325 {
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2326 Lisp_Object before, after;
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2327
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2328 before = Fsubstring (string, make_number (0),
12807
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2329 make_number (search_regs.start[sub]));
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2330 after = Fsubstring (string, make_number (search_regs.end[sub]), Qnil);
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2331
17225
739e41eed8b6 (Freplace_match): Give error if
Richard M. Stallman <rms@gnu.org>
parents: 17102
diff changeset
2332 /* Substitute parts of the match into NEWTEXT
739e41eed8b6 (Freplace_match): Give error if
Richard M. Stallman <rms@gnu.org>
parents: 17102
diff changeset
2333 if desired. */
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2334 if (NILP (literal))
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2335 {
21988
8cf3bbc89c3c (Freplace_match): Fix the loop for copying text
Richard M. Stallman <rms@gnu.org>
parents: 21945
diff changeset
2336 int lastpos = 0;
8cf3bbc89c3c (Freplace_match): Fix the loop for copying text
Richard M. Stallman <rms@gnu.org>
parents: 21945
diff changeset
2337 int lastpos_byte = 0;
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2338 /* We build up the substituted string in ACCUM. */
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2339 Lisp_Object accum;
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2340 Lisp_Object middle;
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2341 int length = STRING_BYTES (XSTRING (newtext));
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2342
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2343 accum = Qnil;
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2344
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2345 for (pos_byte = 0, pos = 0; pos_byte < length;)
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2346 {
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2347 int substart = -1;
31829
43566b0aec59 Avoid some more compiler warnings.
Gerd Moellmann <gerd@gnu.org>
parents: 31486
diff changeset
2348 int subend = 0;
12148
a1c38b9b0f73 (Freplace_match): Do the right thing with backslash.
Karl Heuer <kwzh@gnu.org>
parents: 12147
diff changeset
2349 int delbackslash = 0;
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2350
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
2351 FETCH_STRING_CHAR_ADVANCE (c, newtext, pos, pos_byte);
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
2352
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2353 if (c == '\\')
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2354 {
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
2355 FETCH_STRING_CHAR_ADVANCE (c, newtext, pos, pos_byte);
29018
2f43c508a9b5 (wordify): Use FETCH_STRING_CHAR_ADVANCE
Kenichi Handa <handa@m17n.org>
parents: 28886
diff changeset
2356
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2357 if (c == '&')
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2358 {
12807
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2359 substart = search_regs.start[sub];
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2360 subend = search_regs.end[sub];
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2361 }
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2362 else if (c >= '1' && c <= '9' && c <= search_regs.num_regs + '0')
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2363 {
12147
9200a0e153d3 (Freplace_match): Fix check for valid reg in string replace.
Karl Heuer <kwzh@gnu.org>
parents: 12092
diff changeset
2364 if (search_regs.start[c - '0'] >= 0)
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2365 {
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2366 substart = search_regs.start[c - '0'];
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2367 subend = search_regs.end[c - '0'];
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2368 }
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2369 }
12148
a1c38b9b0f73 (Freplace_match): Do the right thing with backslash.
Karl Heuer <kwzh@gnu.org>
parents: 12147
diff changeset
2370 else if (c == '\\')
a1c38b9b0f73 (Freplace_match): Do the right thing with backslash.
Karl Heuer <kwzh@gnu.org>
parents: 12147
diff changeset
2371 delbackslash = 1;
17225
739e41eed8b6 (Freplace_match): Give error if
Richard M. Stallman <rms@gnu.org>
parents: 17102
diff changeset
2372 else
739e41eed8b6 (Freplace_match): Give error if
Richard M. Stallman <rms@gnu.org>
parents: 17102
diff changeset
2373 error ("Invalid use of `\\' in replacement text");
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2374 }
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2375 if (substart >= 0)
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2376 {
21988
8cf3bbc89c3c (Freplace_match): Fix the loop for copying text
Richard M. Stallman <rms@gnu.org>
parents: 21945
diff changeset
2377 if (pos - 2 != lastpos)
8cf3bbc89c3c (Freplace_match): Fix the loop for copying text
Richard M. Stallman <rms@gnu.org>
parents: 21945
diff changeset
2378 middle = substring_both (newtext, lastpos,
8cf3bbc89c3c (Freplace_match): Fix the loop for copying text
Richard M. Stallman <rms@gnu.org>
parents: 21945
diff changeset
2379 lastpos_byte,
8cf3bbc89c3c (Freplace_match): Fix the loop for copying text
Richard M. Stallman <rms@gnu.org>
parents: 21945
diff changeset
2380 pos - 2, pos_byte - 2);
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2381 else
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2382 middle = Qnil;
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2383 accum = concat3 (accum, middle,
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
2384 Fsubstring (string,
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
2385 make_number (substart),
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2386 make_number (subend)));
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2387 lastpos = pos;
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
2388 lastpos_byte = pos_byte;
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2389 }
12148
a1c38b9b0f73 (Freplace_match): Do the right thing with backslash.
Karl Heuer <kwzh@gnu.org>
parents: 12147
diff changeset
2390 else if (delbackslash)
a1c38b9b0f73 (Freplace_match): Do the right thing with backslash.
Karl Heuer <kwzh@gnu.org>
parents: 12147
diff changeset
2391 {
21988
8cf3bbc89c3c (Freplace_match): Fix the loop for copying text
Richard M. Stallman <rms@gnu.org>
parents: 21945
diff changeset
2392 middle = substring_both (newtext, lastpos,
8cf3bbc89c3c (Freplace_match): Fix the loop for copying text
Richard M. Stallman <rms@gnu.org>
parents: 21945
diff changeset
2393 lastpos_byte,
8cf3bbc89c3c (Freplace_match): Fix the loop for copying text
Richard M. Stallman <rms@gnu.org>
parents: 21945
diff changeset
2394 pos - 1, pos_byte - 1);
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
2395
12148
a1c38b9b0f73 (Freplace_match): Do the right thing with backslash.
Karl Heuer <kwzh@gnu.org>
parents: 12147
diff changeset
2396 accum = concat2 (accum, middle);
a1c38b9b0f73 (Freplace_match): Do the right thing with backslash.
Karl Heuer <kwzh@gnu.org>
parents: 12147
diff changeset
2397 lastpos = pos;
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
2398 lastpos_byte = pos_byte;
12148
a1c38b9b0f73 (Freplace_match): Do the right thing with backslash.
Karl Heuer <kwzh@gnu.org>
parents: 12147
diff changeset
2399 }
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2400 }
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2401
21988
8cf3bbc89c3c (Freplace_match): Fix the loop for copying text
Richard M. Stallman <rms@gnu.org>
parents: 21945
diff changeset
2402 if (pos != lastpos)
8cf3bbc89c3c (Freplace_match): Fix the loop for copying text
Richard M. Stallman <rms@gnu.org>
parents: 21945
diff changeset
2403 middle = substring_both (newtext, lastpos,
8cf3bbc89c3c (Freplace_match): Fix the loop for copying text
Richard M. Stallman <rms@gnu.org>
parents: 21945
diff changeset
2404 lastpos_byte,
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
2405 pos, pos_byte);
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2406 else
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2407 middle = Qnil;
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2408
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2409 newtext = concat2 (accum, middle);
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2410 }
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2411
17225
739e41eed8b6 (Freplace_match): Give error if
Richard M. Stallman <rms@gnu.org>
parents: 17102
diff changeset
2412 /* Do case substitution in NEWTEXT if desired. */
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2413 if (case_action == all_caps)
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2414 newtext = Fupcase (newtext);
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2415 else if (case_action == cap_initial)
12092
b932b2ed40f5 (Freplace_match): Calls to upcase_initials and upcase_initials_region changed
Karl Heuer <kwzh@gnu.org>
parents: 12069
diff changeset
2416 newtext = Fupcase_initials (newtext);
9029
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2417
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2418 return concat3 (before, newtext, after);
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2419 }
f0d89b62dd27 (Freplace_match): New 4th arg OBJECT can specify string to replace in.
Richard M. Stallman <rms@gnu.org>
parents: 8950
diff changeset
2420
18077
27a0ced43e7e (Freplace_match): Use move_if_not_intangible
Richard M. Stallman <rms@gnu.org>
parents: 17463
diff changeset
2421 /* Record point, the move (quietly) to the start of the match. */
23790
c1dbb92db43e (Freplace_match): Set OPOINT clearly for the case
Richard M. Stallman <rms@gnu.org>
parents: 22640
diff changeset
2422 if (PT >= search_regs.end[sub])
18077
27a0ced43e7e (Freplace_match): Use move_if_not_intangible
Richard M. Stallman <rms@gnu.org>
parents: 17463
diff changeset
2423 opoint = PT - ZV;
23790
c1dbb92db43e (Freplace_match): Set OPOINT clearly for the case
Richard M. Stallman <rms@gnu.org>
parents: 22640
diff changeset
2424 else if (PT > search_regs.start[sub])
c1dbb92db43e (Freplace_match): Set OPOINT clearly for the case
Richard M. Stallman <rms@gnu.org>
parents: 22640
diff changeset
2425 opoint = search_regs.end[sub] - ZV;
18077
27a0ced43e7e (Freplace_match): Use move_if_not_intangible
Richard M. Stallman <rms@gnu.org>
parents: 17463
diff changeset
2426 else
27a0ced43e7e (Freplace_match): Use move_if_not_intangible
Richard M. Stallman <rms@gnu.org>
parents: 17463
diff changeset
2427 opoint = PT;
27a0ced43e7e (Freplace_match): Use move_if_not_intangible
Richard M. Stallman <rms@gnu.org>
parents: 17463
diff changeset
2428
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
2429 TEMP_SET_PT (search_regs.start[sub]);
18077
27a0ced43e7e (Freplace_match): Use move_if_not_intangible
Richard M. Stallman <rms@gnu.org>
parents: 17463
diff changeset
2430
2655
594a33ffed85 * search.c (Freplace_match): Arrange for markers sitting at the
Jim Blandy <jimb@redhat.com>
parents: 2475
diff changeset
2431 /* We insert the replacement text before the old text, and then
594a33ffed85 * search.c (Freplace_match): Arrange for markers sitting at the
Jim Blandy <jimb@redhat.com>
parents: 2475
diff changeset
2432 delete the original text. This means that markers at the
594a33ffed85 * search.c (Freplace_match): Arrange for markers sitting at the
Jim Blandy <jimb@redhat.com>
parents: 2475
diff changeset
2433 beginning or end of the original will float to the corresponding
594a33ffed85 * search.c (Freplace_match): Arrange for markers sitting at the
Jim Blandy <jimb@redhat.com>
parents: 2475
diff changeset
2434 position in the replacement. */
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2435 if (!NILP (literal))
4882
8c09f87f5087 (Freplace_match): Fix argument names to match doc string.
Brian Fox <bfox@gnu.org>
parents: 4832
diff changeset
2436 Finsert_and_inherit (1, &newtext);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2437 else
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2438 {
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2439 int length = STRING_BYTES (XSTRING (newtext));
26982
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2440 unsigned char *substed;
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2441 int substed_alloc_size, substed_len;
28387
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2442 int buf_multibyte = !NILP (current_buffer->enable_multibyte_characters);
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2443 int str_multibyte = STRING_MULTIBYTE (newtext);
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2444 Lisp_Object rev_tbl;
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2445
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2446 rev_tbl= (!buf_multibyte && CHAR_TABLE_P (Vnonascii_translation_table)
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2447 ? Fchar_table_extra_slot (Vnonascii_translation_table,
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2448 make_number (0))
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2449 : Qnil);
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2450
26982
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2451 substed_alloc_size = length * 2 + 100;
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2452 substed = (unsigned char *) xmalloc (substed_alloc_size + 1);
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2453 substed_len = 0;
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2454
28387
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2455 /* Go thru NEWTEXT, producing the actual text to insert in
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2456 SUBSTED while adjusting multibyteness to that of the current
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2457 buffer. */
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2458
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2459 for (pos_byte = 0, pos = 0; pos_byte < length;)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2460 {
26982
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2461 unsigned char str[MAX_MULTIBYTE_LENGTH];
28886
3f60536745bd (Freplace_match): Handle case of `\N' in the
Gerd Moellmann <gerd@gnu.org>
parents: 28507
diff changeset
2462 unsigned char *add_stuff = NULL;
3f60536745bd (Freplace_match): Handle case of `\N' in the
Gerd Moellmann <gerd@gnu.org>
parents: 28507
diff changeset
2463 int add_len = 0;
26982
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2464 int idx = -1;
2655
594a33ffed85 * search.c (Freplace_match): Arrange for markers sitting at the
Jim Blandy <jimb@redhat.com>
parents: 2475
diff changeset
2465
28387
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2466 if (str_multibyte)
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2467 {
29018
2f43c508a9b5 (wordify): Use FETCH_STRING_CHAR_ADVANCE
Kenichi Handa <handa@m17n.org>
parents: 28886
diff changeset
2468 FETCH_STRING_CHAR_ADVANCE_NO_CHECK (c, newtext, pos, pos_byte);
28387
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2469 if (!buf_multibyte)
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2470 c = multibyte_char_to_unibyte (c, rev_tbl);
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2471 }
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2472 else
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2473 {
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2474 /* Note that we don't have to increment POS. */
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2475 c = XSTRING (newtext)->data[pos_byte++];
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2476 if (buf_multibyte)
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2477 c = unibyte_char_to_multibyte (c);
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2478 }
22533
6eae236a4f01 (Freplace_match): Work by chars, not by bytes,
Karl Heuer <kwzh@gnu.org>
parents: 22221
diff changeset
2479
26982
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2480 /* Either set ADD_STUFF and ADD_LEN to the text to put in SUBSTED,
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2481 or set IDX to a match index, which means put that part
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2482 of the buffer text into SUBSTED. */
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2483
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2484 if (c == '\\')
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2485 {
28387
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2486 if (str_multibyte)
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2487 {
29018
2f43c508a9b5 (wordify): Use FETCH_STRING_CHAR_ADVANCE
Kenichi Handa <handa@m17n.org>
parents: 28886
diff changeset
2488 FETCH_STRING_CHAR_ADVANCE_NO_CHECK (c, newtext,
2f43c508a9b5 (wordify): Use FETCH_STRING_CHAR_ADVANCE
Kenichi Handa <handa@m17n.org>
parents: 28886
diff changeset
2489 pos, pos_byte);
28387
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2490 if (!buf_multibyte && !SINGLE_BYTE_CHAR_P (c))
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2491 c = multibyte_char_to_unibyte (c, rev_tbl);
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2492 }
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2493 else
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2494 {
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2495 c = XSTRING (newtext)->data[pos_byte++];
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2496 if (buf_multibyte)
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2497 c = unibyte_char_to_multibyte (c);
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2498 }
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2499
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2500 if (c == '&')
26982
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2501 idx = sub;
7856
9687141f6264 (Freplace_match): Be sure not to treat non-digit like digit.
Richard M. Stallman <rms@gnu.org>
parents: 7674
diff changeset
2502 else if (c >= '1' && c <= '9' && c <= search_regs.num_regs + '0')
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2503 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2504 if (search_regs.start[c - '0'] >= 1)
26982
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2505 idx = c - '0';
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2506 }
17225
739e41eed8b6 (Freplace_match): Give error if
Richard M. Stallman <rms@gnu.org>
parents: 17102
diff changeset
2507 else if (c == '\\')
26982
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2508 add_len = 1, add_stuff = "\\";
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2509 else
28387
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2510 {
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2511 xfree (substed);
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2512 error ("Invalid use of `\\' in replacement text");
9a8814cc543c (Freplace_match): Adjust multibyteness of the current
Kenichi Handa <handa@m17n.org>
parents: 27884
diff changeset
2513 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2514 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2515 else
26982
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2516 {
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2517 add_len = CHAR_STRING (c, str);
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2518 add_stuff = str;
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2519 }
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2520
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2521 /* If we want to copy part of a previous match,
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2522 set up ADD_STUFF and ADD_LEN to point to it. */
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2523 if (idx >= 0)
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2524 {
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2525 int begbyte = CHAR_TO_BYTE (search_regs.start[idx]);
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2526 add_len = CHAR_TO_BYTE (search_regs.end[idx]) - begbyte;
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2527 if (search_regs.start[idx] < GPT && GPT < search_regs.end[idx])
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2528 move_gap (search_regs.start[idx]);
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2529 add_stuff = BYTE_POS_ADDR (begbyte);
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2530 }
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2531
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2532 /* Now the stuff we want to add to SUBSTED
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2533 is invariably ADD_LEN bytes starting at ADD_STUFF. */
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2534
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2535 /* Make sure SUBSTED is big enough. */
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2536 if (substed_len + add_len >= substed_alloc_size)
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2537 {
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2538 substed_alloc_size = substed_len + add_len + 500;
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2539 substed = (unsigned char *) xrealloc (substed,
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2540 substed_alloc_size + 1);
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2541 }
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2542
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2543 /* Now add to the end of SUBSTED. */
28886
3f60536745bd (Freplace_match): Handle case of `\N' in the
Gerd Moellmann <gerd@gnu.org>
parents: 28507
diff changeset
2544 if (add_stuff)
3f60536745bd (Freplace_match): Handle case of `\N' in the
Gerd Moellmann <gerd@gnu.org>
parents: 28507
diff changeset
2545 {
3f60536745bd (Freplace_match): Handle case of `\N' in the
Gerd Moellmann <gerd@gnu.org>
parents: 28507
diff changeset
2546 bcopy (add_stuff, substed + substed_len, add_len);
3f60536745bd (Freplace_match): Handle case of `\N' in the
Gerd Moellmann <gerd@gnu.org>
parents: 28507
diff changeset
2547 substed_len += add_len;
3f60536745bd (Freplace_match): Handle case of `\N' in the
Gerd Moellmann <gerd@gnu.org>
parents: 28507
diff changeset
2548 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2549 }
26982
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2550
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2551 /* Now insert what we accumulated. */
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2552 insert_and_inherit (substed, substed_len);
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2553
3527c131b069 (Freplace_match): For nonliteral replacement,
Richard M. Stallman <rms@gnu.org>
parents: 26869
diff changeset
2554 xfree (substed);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2555 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2556
16039
855c8d8ba0f0 Change all references from point to PT.
Karl Heuer <kwzh@gnu.org>
parents: 15667
diff changeset
2557 inslen = PT - (search_regs.start[sub]);
12807
34d269b30df1 (Freplace_match): New arg SUBEXP.
Richard M. Stallman <rms@gnu.org>
parents: 12244
diff changeset
2558 del_range (search_regs.start[sub] + inslen, search_regs.end[sub] + inslen);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2559
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2560 if (case_action == all_caps)
16039
855c8d8ba0f0 Change all references from point to PT.
Karl Heuer <kwzh@gnu.org>
parents: 15667
diff changeset
2561 Fupcase_region (make_number (PT - inslen), make_number (PT));
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2562 else if (case_action == cap_initial)
16039
855c8d8ba0f0 Change all references from point to PT.
Karl Heuer <kwzh@gnu.org>
parents: 15667
diff changeset
2563 Fupcase_initials_region (make_number (PT - inslen), make_number (PT));
18077
27a0ced43e7e (Freplace_match): Use move_if_not_intangible
Richard M. Stallman <rms@gnu.org>
parents: 17463
diff changeset
2564
18081
300068b4fcef (Freplace_match): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 18077
diff changeset
2565 newpoint = PT;
300068b4fcef (Freplace_match): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 18077
diff changeset
2566
18077
27a0ced43e7e (Freplace_match): Use move_if_not_intangible
Richard M. Stallman <rms@gnu.org>
parents: 17463
diff changeset
2567 /* Put point back where it was in the text. */
18124
6f2c80d2425a (Freplace_match): If opoint is 0, that's relative to ZV.
Richard M. Stallman <rms@gnu.org>
parents: 18112
diff changeset
2568 if (opoint <= 0)
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
2569 TEMP_SET_PT (opoint + ZV);
18077
27a0ced43e7e (Freplace_match): Use move_if_not_intangible
Richard M. Stallman <rms@gnu.org>
parents: 17463
diff changeset
2570 else
20545
c20c92ff4055 (looking_at_1): Use bytepos to call re_search_2.
Richard M. Stallman <rms@gnu.org>
parents: 20347
diff changeset
2571 TEMP_SET_PT (opoint);
18077
27a0ced43e7e (Freplace_match): Use move_if_not_intangible
Richard M. Stallman <rms@gnu.org>
parents: 17463
diff changeset
2572
27a0ced43e7e (Freplace_match): Use move_if_not_intangible
Richard M. Stallman <rms@gnu.org>
parents: 17463
diff changeset
2573 /* Now move point "officially" to the start of the inserted replacement. */
18081
300068b4fcef (Freplace_match): Fix previous change.
Richard M. Stallman <rms@gnu.org>
parents: 18077
diff changeset
2574 move_if_not_intangible (newpoint);
18077
27a0ced43e7e (Freplace_match): Use move_if_not_intangible
Richard M. Stallman <rms@gnu.org>
parents: 17463
diff changeset
2575
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2576 return Qnil;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2577 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2578
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2579 static Lisp_Object
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2580 match_limit (num, beginningp)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2581 Lisp_Object num;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2582 int beginningp;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2583 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2584 register int n;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2585
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2586 CHECK_NUMBER (num, 0);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2587 n = XINT (num);
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2588 if (n < 0 || n >= search_regs.num_regs)
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2589 args_out_of_range (num, make_number (search_regs.num_regs));
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2590 if (search_regs.num_regs <= 0
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2591 || search_regs.start[n] < 0)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2592 return Qnil;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2593 return (make_number ((beginningp) ? search_regs.start[n]
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2594 : search_regs.end[n]));
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2595 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2596
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2597 DEFUN ("match-beginning", Fmatch_beginning, Smatch_beginning, 1, 1, 0,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2598 "Return position of start of text matched by last search.\n\
14086
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2599 SUBEXP, a number, specifies which parenthesized expression in the last\n\
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2600 regexp.\n\
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2601 Value is nil if SUBEXPth pair didn't match, or there were less than\n\
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2602 SUBEXP pairs.\n\
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2603 Zero means the entire text matched by the whole regexp or whole string.")
14086
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2604 (subexp)
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2605 Lisp_Object subexp;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2606 {
14086
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2607 return match_limit (subexp, 1);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2608 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2609
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2610 DEFUN ("match-end", Fmatch_end, Smatch_end, 1, 1, 0,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2611 "Return position of end of text matched by last search.\n\
14086
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2612 SUBEXP, a number, specifies which parenthesized expression in the last\n\
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2613 regexp.\n\
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2614 Value is nil if SUBEXPth pair didn't match, or there were less than\n\
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2615 SUBEXP pairs.\n\
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2616 Zero means the entire text matched by the whole regexp or whole string.")
14086
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2617 (subexp)
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2618 Lisp_Object subexp;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2619 {
14086
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2620 return match_limit (subexp, 0);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2621 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2622
16724
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2623 DEFUN ("match-data", Fmatch_data, Smatch_data, 0, 2, 0,
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2624 "Return a list containing all info on what the last search matched.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2625 Element 2N is `(match-beginning N)'; element 2N + 1 is `(match-end N)'.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2626 All the elements are markers or nil (nil if the Nth pair didn't match)\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2627 if the last match was on a buffer; integers or nil if a string was matched.\n\
16724
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2628 Use `store-match-data' to reinstate the data in this list.\n\
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2629 \n\
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2630 If INTEGERS (the optional first argument) is non-nil, always use integers\n\
16731
039ef6e74d3a (Fmatch_data): Doc fix.
Richard M. Stallman <rms@gnu.org>
parents: 16724
diff changeset
2631 \(rather than markers) to represent buffer positions.\n\
16724
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2632 If REUSE is a list, reuse it as part of the value. If REUSE is long enough\n\
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2633 to hold all the values, and if INTEGERS is non-nil, no consing is done.")
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2634 (integers, reuse)
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2635 Lisp_Object integers, reuse;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2636 {
16724
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2637 Lisp_Object tail, prev;
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2638 Lisp_Object *data;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2639 int i, len;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2640
727
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2641 if (NILP (last_thing_searched))
15667
9531c03134b6 (Fmatch_data): If no matching done yet, return Qnil.
Karl Heuer <kwzh@gnu.org>
parents: 14186
diff changeset
2642 return Qnil;
727
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2643
31829
43566b0aec59 Avoid some more compiler warnings.
Gerd Moellmann <gerd@gnu.org>
parents: 31486
diff changeset
2644 prev = Qnil;
43566b0aec59 Avoid some more compiler warnings.
Gerd Moellmann <gerd@gnu.org>
parents: 31486
diff changeset
2645
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2646 data = (Lisp_Object *) alloca ((2 * search_regs.num_regs)
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2647 * sizeof (Lisp_Object));
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2648
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2649 len = -1;
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2650 for (i = 0; i < search_regs.num_regs; i++)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2651 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2652 int start = search_regs.start[i];
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2653 if (start >= 0)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2654 {
16724
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2655 if (EQ (last_thing_searched, Qt)
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2656 || ! NILP (integers))
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2657 {
9319
7969182b6cc6 (skip_chars, Fmatch_data, Fstore_match_data): Don't use XFASTINT as an lvalue.
Karl Heuer <kwzh@gnu.org>
parents: 9278
diff changeset
2658 XSETFASTINT (data[2 * i], start);
7969182b6cc6 (skip_chars, Fmatch_data, Fstore_match_data): Don't use XFASTINT as an lvalue.
Karl Heuer <kwzh@gnu.org>
parents: 9278
diff changeset
2659 XSETFASTINT (data[2 * i + 1], search_regs.end[i]);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2660 }
9113
766b6288e0f2 (Fmatch_data, Fstore_match_data): Use type test macros.
Karl Heuer <kwzh@gnu.org>
parents: 9029
diff changeset
2661 else if (BUFFERP (last_thing_searched))
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2662 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2663 data[2 * i] = Fmake_marker ();
727
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2664 Fset_marker (data[2 * i],
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2665 make_number (start),
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2666 last_thing_searched);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2667 data[2 * i + 1] = Fmake_marker ();
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2668 Fset_marker (data[2 * i + 1],
727
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2669 make_number (search_regs.end[i]),
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2670 last_thing_searched);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2671 }
727
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2672 else
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2673 /* last_thing_searched must always be Qt, a buffer, or Qnil. */
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2674 abort ();
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2675
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2676 len = i;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2677 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2678 else
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2679 data[2 * i] = data [2 * i + 1] = Qnil;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2680 }
16724
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2681
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2682 /* If REUSE is not usable, cons up the values and return them. */
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2683 if (! CONSP (reuse))
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2684 return Flist (2 * len + 2, data);
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2685
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2686 /* If REUSE is a list, store as many value elements as will fit
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2687 into the elements of REUSE. */
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2688 for (i = 0, tail = reuse; CONSP (tail);
25663
a5eaace0fa01 Use XCAR and XCDR instead of explicit member access.
Ken Raeburn <raeburn@raeburn.org>
parents: 25441
diff changeset
2689 i++, tail = XCDR (tail))
16724
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2690 {
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2691 if (i < 2 * len + 2)
25663
a5eaace0fa01 Use XCAR and XCDR instead of explicit member access.
Ken Raeburn <raeburn@raeburn.org>
parents: 25441
diff changeset
2692 XCAR (tail) = data[i];
16724
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2693 else
25663
a5eaace0fa01 Use XCAR and XCDR instead of explicit member access.
Ken Raeburn <raeburn@raeburn.org>
parents: 25441
diff changeset
2694 XCAR (tail) = Qnil;
16724
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2695 prev = tail;
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2696 }
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2697
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2698 /* If we couldn't fit all value elements into REUSE,
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2699 cons up the rest of them and add them to the end of REUSE. */
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2700 if (i < 2 * len + 2)
25663
a5eaace0fa01 Use XCAR and XCDR instead of explicit member access.
Ken Raeburn <raeburn@raeburn.org>
parents: 25441
diff changeset
2701 XCDR (prev) = Flist (2 * len + 2 - i, data + i);
16724
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2702
4b1fb372a4fe (Fmatch_data): New args INTEGERS and REUSE.
Richard M. Stallman <rms@gnu.org>
parents: 16275
diff changeset
2703 return reuse;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2704 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2705
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2706
21171
60f6085df198 (Fset_match_data): Renamed from Fstore_match_data.
Richard M. Stallman <rms@gnu.org>
parents: 21117
diff changeset
2707 DEFUN ("set-match-data", Fset_match_data, Sset_match_data, 1, 1, 0,
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2708 "Set internal data on last search match from elements of LIST.\n\
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2709 LIST should have been created by calling `match-data' previously.")
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2710 (list)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2711 register Lisp_Object list;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2712 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2713 register int i;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2714 register Lisp_Object marker;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2715
10032
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2716 if (running_asynch_code)
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2717 save_search_regs ();
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2718
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2719 if (!CONSP (list) && !NILP (list))
1926
952f2a18f83d * callint.c (Fcall_interactively): Pass the correct number of
Jim Blandy <jimb@redhat.com>
parents: 1896
diff changeset
2720 list = wrong_type_argument (Qconsp, list);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2721
727
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2722 /* Unless we find a marker with a buffer in LIST, assume that this
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2723 match data came from a string. */
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2724 last_thing_searched = Qt;
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2725
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2726 /* Allocate registers if they don't already exist. */
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2727 {
1523
bd61aaa7828b * search.c (Fstore_match_data): Don't assume Flength returns an
Jim Blandy <jimb@redhat.com>
parents: 1413
diff changeset
2728 int length = XFASTINT (Flength (list)) / 2;
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2729
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2730 if (length > search_regs.num_regs)
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2731 {
708
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2732 if (search_regs.num_regs == 0)
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2733 {
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2734 search_regs.start
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2735 = (regoff_t *) xmalloc (length * sizeof (regoff_t));
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2736 search_regs.end
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2737 = (regoff_t *) xmalloc (length * sizeof (regoff_t));
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2738 }
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2739 else
708
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2740 {
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2741 search_regs.start
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2742 = (regoff_t *) xrealloc (search_regs.start,
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2743 length * sizeof (regoff_t));
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2744 search_regs.end
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2745 = (regoff_t *) xrealloc (search_regs.end,
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2746 length * sizeof (regoff_t));
030fb4635335 *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 648
diff changeset
2747 }
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2748
33052
9ec478daa468 (Fset_match_data): Be sure to make search_regs always sane.
Kenichi Handa <handa@m17n.org>
parents: 32387
diff changeset
2749 for (i = search_regs.num_regs; i < length; i++)
9ec478daa468 (Fset_match_data): Be sure to make search_regs always sane.
Kenichi Handa <handa@m17n.org>
parents: 32387
diff changeset
2750 search_regs.start[i] = -1;
9ec478daa468 (Fset_match_data): Be sure to make search_regs always sane.
Kenichi Handa <handa@m17n.org>
parents: 32387
diff changeset
2751
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
2752 search_regs.num_regs = length;
621
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2753 }
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2754 }
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2755
eca8812e61cd *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 605
diff changeset
2756 for (i = 0; i < search_regs.num_regs; i++)
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2757 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2758 marker = Fcar (list);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2759 if (NILP (marker))
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2760 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2761 search_regs.start[i] = -1;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2762 list = Fcdr (list);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2763 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2764 else
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2765 {
33052
9ec478daa468 (Fset_match_data): Be sure to make search_regs always sane.
Kenichi Handa <handa@m17n.org>
parents: 32387
diff changeset
2766 int from;
9ec478daa468 (Fset_match_data): Be sure to make search_regs always sane.
Kenichi Handa <handa@m17n.org>
parents: 32387
diff changeset
2767
9113
766b6288e0f2 (Fmatch_data, Fstore_match_data): Use type test macros.
Karl Heuer <kwzh@gnu.org>
parents: 9029
diff changeset
2768 if (MARKERP (marker))
727
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2769 {
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2770 if (XMARKER (marker)->buffer == 0)
9319
7969182b6cc6 (skip_chars, Fmatch_data, Fstore_match_data): Don't use XFASTINT as an lvalue.
Karl Heuer <kwzh@gnu.org>
parents: 9278
diff changeset
2771 XSETFASTINT (marker, 0);
727
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2772 else
9278
f2138d548313 (Flooking_at, skip_chars, search_buffer, set_search_regs, Fstore_match_data):
Karl Heuer <kwzh@gnu.org>
parents: 9113
diff changeset
2773 XSETBUFFER (last_thing_searched, XMARKER (marker)->buffer);
727
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2774 }
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2775
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2776 CHECK_NUMBER_COERCE_MARKER (marker, 0);
33052
9ec478daa468 (Fset_match_data): Be sure to make search_regs always sane.
Kenichi Handa <handa@m17n.org>
parents: 32387
diff changeset
2777 from = XINT (marker);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2778 list = Fcdr (list);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2779
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2780 marker = Fcar (list);
9113
766b6288e0f2 (Fmatch_data, Fstore_match_data): Use type test macros.
Karl Heuer <kwzh@gnu.org>
parents: 9029
diff changeset
2781 if (MARKERP (marker) && XMARKER (marker)->buffer == 0)
9319
7969182b6cc6 (skip_chars, Fmatch_data, Fstore_match_data): Don't use XFASTINT as an lvalue.
Karl Heuer <kwzh@gnu.org>
parents: 9278
diff changeset
2782 XSETFASTINT (marker, 0);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2783
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2784 CHECK_NUMBER_COERCE_MARKER (marker, 0);
33052
9ec478daa468 (Fset_match_data): Be sure to make search_regs always sane.
Kenichi Handa <handa@m17n.org>
parents: 32387
diff changeset
2785 search_regs.start[i] = from;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2786 search_regs.end[i] = XINT (marker);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2787 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2788 list = Fcdr (list);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2789 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2790
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2791 return Qnil;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2792 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2793
10032
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2794 /* If non-zero the match data have been saved in saved_search_regs
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2795 during the execution of a sentinel or filter. */
10128
59ccd063e016 (search_regs_saved): Delete initializer.
Richard M. Stallman <rms@gnu.org>
parents: 10055
diff changeset
2796 static int search_regs_saved;
10032
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2797 static struct re_registers saved_search_regs;
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2798
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2799 /* Called from Flooking_at, Fstring_match, search_buffer, Fstore_match_data
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2800 if asynchronous code (filter or sentinel) is running. */
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2801 static void
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2802 save_search_regs ()
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2803 {
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2804 if (!search_regs_saved)
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2805 {
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2806 saved_search_regs.num_regs = search_regs.num_regs;
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2807 saved_search_regs.start = search_regs.start;
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2808 saved_search_regs.end = search_regs.end;
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2809 search_regs.num_regs = 0;
10250
422c3b96efda (set_search_regs): Really set search_regs.start and .end.
Richard M. Stallman <rms@gnu.org>
parents: 10141
diff changeset
2810 search_regs.start = 0;
422c3b96efda (set_search_regs): Really set search_regs.start and .end.
Richard M. Stallman <rms@gnu.org>
parents: 10141
diff changeset
2811 search_regs.end = 0;
10032
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2812
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2813 search_regs_saved = 1;
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2814 }
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2815 }
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2816
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2817 /* Called upon exit from filters and sentinels. */
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2818 void
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2819 restore_match_data ()
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2820 {
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2821 if (search_regs_saved)
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2822 {
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2823 if (search_regs.num_regs > 0)
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2824 {
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2825 xfree (search_regs.start);
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2826 xfree (search_regs.end);
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2827 }
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2828 search_regs.num_regs = saved_search_regs.num_regs;
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2829 search_regs.start = saved_search_regs.start;
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2830 search_regs.end = saved_search_regs.end;
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2831
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2832 search_regs_saved = 0;
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2833 }
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2834 }
f689803caa92 Added code for automatically saving and restoring the match data
Francesco Potortì <pot@gnu.org>
parents: 10020
diff changeset
2835
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2836 /* Quote a string to inactivate reg-expr chars */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2837
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2838 DEFUN ("regexp-quote", Fregexp_quote, Sregexp_quote, 1, 1, 0,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2839 "Return a regexp string which matches exactly STRING and nothing else.")
14086
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2840 (string)
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2841 Lisp_Object string;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2842 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2843 register unsigned char *in, *out, *end;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2844 register unsigned char *temp;
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
2845 int backslashes_added = 0;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2846
14086
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2847 CHECK_STRING (string, 0);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2848
21244
50929073a0ba Use STRING_BYTES and SET_STRING_BYTES.
Richard M. Stallman <rms@gnu.org>
parents: 21171
diff changeset
2849 temp = (unsigned char *) alloca (STRING_BYTES (XSTRING (string)) * 2);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2850
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2851 /* Now copy the data into the new string, inserting escapes. */
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2852
14086
a410808fda15 (Fmatch_end, Fregexp_quote): Harmonize arguments with documentation.
Erik Naggum <erik@naggum.no>
parents: 14036
diff changeset
2853 in = XSTRING (string)->data;
21244
50929073a0ba Use STRING_BYTES and SET_STRING_BYTES.
Richard M. Stallman <rms@gnu.org>
parents: 21171
diff changeset
2854 end = in + STRING_BYTES (XSTRING (string));
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2855 out = temp;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2856
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2857 for (; in != end; in++)
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2858 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2859 if (*in == '[' || *in == ']'
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2860 || *in == '*' || *in == '.' || *in == '\\'
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2861 || *in == '?' || *in == '+'
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2862 || *in == '^' || *in == '$')
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
2863 *out++ = '\\', backslashes_added++;
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2864 *out++ = *in;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2865 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2866
21248
aba5e1c3328b (Fregexp_quote): Use make_specified_string.
Richard M. Stallman <rms@gnu.org>
parents: 21244
diff changeset
2867 return make_specified_string (temp,
20588
138c95482e6b (search_buffer): Handle bytes vs chars in non-RE case.
Richard M. Stallman <rms@gnu.org>
parents: 20547
diff changeset
2868 XSTRING (string)->size + backslashes_added,
21248
aba5e1c3328b (Fregexp_quote): Use make_specified_string.
Richard M. Stallman <rms@gnu.org>
parents: 21244
diff changeset
2869 out - temp,
aba5e1c3328b (Fregexp_quote): Use make_specified_string.
Richard M. Stallman <rms@gnu.org>
parents: 21244
diff changeset
2870 STRING_MULTIBYTE (string));
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2871 }
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2872
21514
fa9ff387d260 Fix -Wimplicit warnings.
Andreas Schwab <schwab@suse.de>
parents: 21457
diff changeset
2873 void
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2874 syms_of_search ()
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2875 {
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2876 register int i;
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2877
9605
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
2878 for (i = 0; i < REGEXP_CACHE_SIZE; ++i)
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
2879 {
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
2880 searchbufs[i].buf.allocated = 100;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
2881 searchbufs[i].buf.buffer = (unsigned char *) malloc (100);
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
2882 searchbufs[i].buf.fastmap = searchbufs[i].fastmap;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
2883 searchbufs[i].regexp = Qnil;
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
2884 staticpro (&searchbufs[i].regexp);
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
2885 searchbufs[i].next = (i == REGEXP_CACHE_SIZE-1 ? 0 : &searchbufs[i+1]);
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
2886 }
cf97e75d8e02 (searchbufs): New variable, replaces searchbuf and last_regexp and
Karl Heuer <kwzh@gnu.org>
parents: 9452
diff changeset
2887 searchbuf_head = &searchbufs[0];
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2888
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2889 Qsearch_failed = intern ("search-failed");
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2890 staticpro (&Qsearch_failed);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2891 Qinvalid_regexp = intern ("invalid-regexp");
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2892 staticpro (&Qinvalid_regexp);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2893
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2894 Fput (Qsearch_failed, Qerror_conditions,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2895 Fcons (Qsearch_failed, Fcons (Qerror, Qnil)));
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2896 Fput (Qsearch_failed, Qerror_message,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2897 build_string ("Search failed"));
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2898
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2899 Fput (Qinvalid_regexp, Qerror_conditions,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2900 Fcons (Qinvalid_regexp, Fcons (Qerror, Qnil)));
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2901 Fput (Qinvalid_regexp, Qerror_message,
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2902 build_string ("Invalid regexp"));
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2903
727
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2904 last_thing_searched = Qnil;
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2905 staticpro (&last_thing_searched);
540b047ece4d *** empty log message ***
Jim Blandy <jimb@redhat.com>
parents: 708
diff changeset
2906
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2907 defsubr (&Slooking_at);
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2908 defsubr (&Sposix_looking_at);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2909 defsubr (&Sstring_match);
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2910 defsubr (&Sposix_string_match);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2911 defsubr (&Ssearch_forward);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2912 defsubr (&Ssearch_backward);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2913 defsubr (&Sword_search_forward);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2914 defsubr (&Sword_search_backward);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2915 defsubr (&Sre_search_forward);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2916 defsubr (&Sre_search_backward);
10020
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2917 defsubr (&Sposix_search_forward);
c41ce96785a8 (struct regexp_cache): New field `posix'.
Richard M. Stallman <rms@gnu.org>
parents: 9605
diff changeset
2918 defsubr (&Sposix_search_backward);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2919 defsubr (&Sreplace_match);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2920 defsubr (&Smatch_beginning);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2921 defsubr (&Smatch_end);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2922 defsubr (&Smatch_data);
21171
60f6085df198 (Fset_match_data): Renamed from Fstore_match_data.
Richard M. Stallman <rms@gnu.org>
parents: 21117
diff changeset
2923 defsubr (&Sset_match_data);
603
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2924 defsubr (&Sregexp_quote);
470f556a9453 Initial revision
Jim Blandy <jimb@redhat.com>
parents:
diff changeset
2925 }