94
88
_sub_named = Replacer()
95
_sub_named.add(r'\[:digit:\]', r'\d')
96
_sub_named.add(r'\[:space:\]', r'\s')
97
_sub_named.add(r'\[:alnum:\]', r'\w')
98
_sub_named.add(r'\[:ascii:\]', r'\0-\x7f')
99
_sub_named.add(r'\[:blank:\]', r' \t')
100
_sub_named.add(r'\[:cntrl:\]', r'\0-\x1f\x7f-\x9f')
89
_sub_named.add(ur'\[:digit:\]', ur'\d')
90
_sub_named.add(ur'\[:space:\]', ur'\s')
91
_sub_named.add(ur'\[:alnum:\]', ur'\w')
92
_sub_named.add(ur'\[:ascii:\]', ur'\0-\x7f')
93
_sub_named.add(ur'\[:blank:\]', ur' \t')
94
_sub_named.add(ur'\[:cntrl:\]', ur'\0-\x1f\x7f-\x9f')
103
97
def _sub_group(m):
130
124
_sub_re = Replacer()
131
125
_sub_re.add(u'^RE:', u'')
132
_sub_re.add(u'\\((?!\\?)', u'(?:')
133
_sub_re.add(u'\\(\\?P<.*>', _invalid_regex(u'(?:'))
134
_sub_re.add(u'\\(\\?P=[^)]*\\)', _invalid_regex(u''))
135
_sub_re.add(r'\\+$', _trailing_backslashes_regex)
126
_sub_re.add(u'\((?!\?)', u'(?:')
127
_sub_re.add(u'\(\?P<.*>', _invalid_regex(u'(?:'))
128
_sub_re.add(u'\(\?P=[^)]*\)', _invalid_regex(u''))
129
_sub_re.add(ur'\\+$', _trailing_backslashes_regex)
138
132
_sub_fullpath = Replacer()
139
_sub_fullpath.add(r'^RE:.*', _sub_re) # RE:<anything> is a regex
140
_sub_fullpath.add(r'\[\^?\]?(?:[^][]|\[:[^]]+:\])+\]',
141
_sub_group) # char group
142
_sub_fullpath.add(r'(?:(?<=/)|^)(?:\.?/)+', u'') # canonicalize path
143
_sub_fullpath.add(r'\\.', r'\&') # keep anything backslashed
144
_sub_fullpath.add(r'[(){}|^$+.]', r'\\&') # escape specials
145
_sub_fullpath.add(r'(?:(?<=/)|^)\*\*+/', r'(?:.*/)?') # **/ after ^ or /
146
_sub_fullpath.add(r'\*+', r'[^/]*') # * elsewhere
147
_sub_fullpath.add(r'\?', r'[^/]') # ? everywhere
133
_sub_fullpath.add(ur'^RE:.*', _sub_re) # RE:<anything> is a regex
134
_sub_fullpath.add(ur'\[\^?\]?(?:[^][]|\[:[^]]+:\])+\]', _sub_group) # char group
135
_sub_fullpath.add(ur'(?:(?<=/)|^)(?:\.?/)+', u'') # canonicalize path
136
_sub_fullpath.add(ur'\\.', ur'\&') # keep anything backslashed
137
_sub_fullpath.add(ur'[(){}|^$+.]', ur'\\&') # escape specials
138
_sub_fullpath.add(ur'(?:(?<=/)|^)\*\*+/', ur'(?:.*/)?') # **/ after ^ or /
139
_sub_fullpath.add(ur'\*+', ur'[^/]*') # * elsewhere
140
_sub_fullpath.add(ur'\?', ur'[^/]') # ? everywhere
150
143
_sub_basename = Replacer()
151
_sub_basename.add(r'\[\^?\]?(?:[^][]|\[:[^]]+:\])+\]',
152
_sub_group) # char group
153
_sub_basename.add(r'\\.', r'\&') # keep anything backslashed
154
_sub_basename.add(r'[(){}|^$+.]', r'\\&') # escape specials
155
_sub_basename.add(r'\*+', r'.*') # * everywhere
156
_sub_basename.add(r'\?', r'.') # ? everywhere
144
_sub_basename.add(ur'\[\^?\]?(?:[^][]|\[:[^]]+:\])+\]', _sub_group) # char group
145
_sub_basename.add(ur'\\.', ur'\&') # keep anything backslashed
146
_sub_basename.add(ur'[(){}|^$+.]', ur'\\&') # escape specials
147
_sub_basename.add(ur'\*+', ur'.*') # * everywhere
148
_sub_basename.add(ur'\?', ur'.') # ? everywhere
159
151
def _sub_extension(pattern):
185
177
so are matched first, then the basename patterns, then the fullpath
188
# We want to _add_patterns in a specific order (as per type_list below)
189
# starting with the shortest and going to the longest.
190
# As some Python version don't support ordered dicts the list below is
191
# used to select inputs for _add_pattern in a specific order.
192
pattern_types = ["extension", "basename", "fullpath"]
196
"translator": _sub_extension,
197
"prefix": r'(?:.*/)?(?!.*/)(?:.*\.)'
200
"translator": _sub_basename,
201
"prefix": r'(?:.*/)?(?!.*/)'
204
"translator": _sub_fullpath,
209
180
def __init__(self, patterns):
210
181
self._regex_patterns = []
216
185
for pat in patterns:
217
186
pat = normalize_pattern(pat)
218
pattern_lists[Globster.identify(pat)].append(pat)
219
pi = Globster.pattern_info
220
for t in Globster.pattern_types:
221
self._add_patterns(pattern_lists[t], pi[t]["translator"],
187
if pat.startswith(u'RE:') or u'/' in pat:
188
path_patterns.append(pat)
189
elif pat.startswith(u'*.'):
190
ext_patterns.append(pat)
192
base_patterns.append(pat)
193
self._add_patterns(ext_patterns,_sub_extension,
194
prefix=r'(?:.*/)?(?!.*/)(?:.*\.)')
195
self._add_patterns(base_patterns,_sub_basename,
196
prefix=r'(?:.*/)?(?!.*/)')
197
self._add_patterns(path_patterns,_sub_fullpath)
224
199
def _add_patterns(self, patterns, translator, prefix=''):
227
'(%s)' % translator(pat) for pat in patterns[:99]]
201
grouped_rules = ['(%s)' % translator(pat) for pat in patterns[:99]]
228
202
joined_rule = '%s(?:%s)$' % (prefix, '|'.join(grouped_rules))
229
# Explicitly use lazy_compile here, because we count on its
230
# nicer error reporting.
231
self._regex_patterns.append((
232
lazy_regex.lazy_compile(joined_rule, re.UNICODE),
203
self._regex_patterns.append((re.compile(joined_rule, re.UNICODE),
234
205
patterns = patterns[99:]
239
210
:return A matching pattern or None if there is no matching pattern.
242
for regex, patterns in self._regex_patterns:
243
match = regex.match(filename)
245
return patterns[match.lastindex - 1]
246
except lazy_regex.InvalidPattern as e:
247
# We can't show the default e.msg to the user as thats for
248
# the combined pattern we sent to regex. Instead we indicate to
249
# the user that an ignore file needs fixing.
250
mutter('Invalid pattern found in regex: %s.', e.msg)
252
"File ~/.config/breezy/ignore or "
253
".bzrignore contains error(s).")
255
for _, patterns in self._regex_patterns:
257
if not Globster.is_pattern_valid(p):
258
bad_patterns += ('\n %s' % p)
259
e.msg += bad_patterns
212
for regex, patterns in self._regex_patterns:
213
match = regex.match(filename)
215
return patterns[match.lastindex -1]
264
def identify(pattern):
265
"""Returns pattern category.
267
:param pattern: normalized pattern.
268
Identify if a pattern is fullpath, basename or extension
269
and returns the appropriate type.
271
if pattern.startswith(u'RE:') or u'/' in pattern:
273
elif pattern.startswith(u'*.'):
279
def is_pattern_valid(pattern):
280
"""Returns True if pattern is valid.
282
:param pattern: Normalized pattern.
283
is_pattern_valid() assumes pattern to be normalized.
284
see: globbing.normalize_pattern
287
translator = Globster.pattern_info[Globster.identify(
288
pattern)]["translator"]
289
tpattern = '(%s)' % translator(pattern)
291
re_obj = lazy_regex.lazy_compile(tpattern, re.UNICODE)
292
re_obj.search("") # force compile
293
except lazy_regex.InvalidPattern:
298
class ExceptionGlobster(object):
299
"""A Globster that supports exception patterns.
301
Exceptions are ignore patterns prefixed with '!'. Exception
302
patterns take precedence over regular patterns and cause a
303
matching filename to return None from the match() function.
304
Patterns using a '!!' prefix are highest precedence, and act
305
as regular ignores. '!!' patterns are useful to establish ignores
306
that apply under paths specified by '!' exception patterns.
309
def __init__(self, patterns):
310
ignores = [[], [], []]
312
if p.startswith(u'!!'):
313
ignores[2].append(p[2:])
314
elif p.startswith(u'!'):
315
ignores[1].append(p[1:])
318
self._ignores = [Globster(i) for i in ignores]
320
def match(self, filename):
321
"""Searches for a pattern that matches the given filename.
323
:return A matching pattern or None if there is no matching pattern.
325
double_neg = self._ignores[2].match(filename)
327
return "!!%s" % double_neg
328
elif self._ignores[1].match(filename):
331
return self._ignores[0].match(filename)
334
219
class _OrderedGlobster(Globster):
335
220
"""A Globster that keeps pattern order."""
343
228
self._regex_patterns = []
344
229
for pat in patterns:
345
230
pat = normalize_pattern(pat)
346
t = Globster.identify(pat)
347
self._add_patterns([pat], Globster.pattern_info[t]["translator"],
348
Globster.pattern_info[t]["prefix"])
351
_slashes = lazy_regex.lazy_compile(r'[\\/]+')
231
if pat.startswith(u'RE:') or u'/' in pat:
232
self._add_patterns([pat], _sub_fullpath)
233
elif pat.startswith(u'*.'):
234
self._add_patterns([pat], _sub_extension,
235
prefix=r'(?:.*/)?(?!.*/)(?:.*\.)')
237
self._add_patterns([pat], _sub_basename,
238
prefix=r'(?:.*/)?(?!.*/)')
354
241
def normalize_pattern(pattern):