upstream/mercurial-mirror Commit - r24448:55c44934

match: add isexact() method to hide internals...

Martin von Zweigbergk -

r24448:55c44934 default

parent child

mercurial/dirstate.py

0 +3 -3

              # dirstate.py - working directory tracking for mercurial
              #
              # Copyright 2005-2007 Matt Mackall <mpm@selenic.com>
              #
              # This software may be used and distributed according to the terms of the
              # GNU General Public License version 2 or any later version.
              from node import nullid
              from i18n import _
              import scmutil, util, ignore, osutil, parsers, encoding, pathutil
              import os, stat, errno
              propertycache = util.propertycache
              filecache = scmutil.filecache
              _rangemask = 0x7fffffff
              dirstatetuple = parsers.dirstatetuple
              class repocache(filecache):
                  """filecache for files in .hg/"""
                  def join(self, obj, fname):
                      return obj._opener.join(fname)
              class rootcache(filecache):
                  """filecache for files in the repository root"""
                  def join(self, obj, fname):
                      return obj._join(fname)
              class dirstate(object):
                  def __init__(self, opener, ui, root, validate):
                      '''Create a new dirstate object.
                      opener is an open()-like callable that can be used to open the
                      dirstate file; root is the root of the directory tracked by
                      the dirstate.
                      '''
                      self._opener = opener
                      self._validate = validate
                      self._root = root
                      # ntpath.join(root, '') of Python 2.7.9 does not add sep if root is
                      # UNC path pointing to root share (issue4557)
                      if root.endswith(os.sep):
                          self._rootdir = root
                      else:
                          self._rootdir = root + os.sep
                      self._dirty = False
                      self._dirtypl = False
                      self._lastnormaltime = 0
                      self._ui = ui
                      self._filecache = {}
                      self._parentwriters = 0
                  def beginparentchange(self):
                      '''Marks the beginning of a set of changes that involve changing
                      the dirstate parents. If there is an exception during this time,
                      the dirstate will not be written when the wlock is released. This
                      prevents writing an incoherent dirstate where the parent doesn't
                      match the contents.
                      '''
                      self._parentwriters += 1
                  def endparentchange(self):
                      '''Marks the end of a set of changes that involve changing the
                      dirstate parents. Once all parent changes have been marked done,
                      the wlock will be free to write the dirstate on release.
                      '''
                      if self._parentwriters > 0:
                          self._parentwriters -= 1
                  def pendingparentchange(self):
                      '''Returns true if the dirstate is in the middle of a set of changes
                      that modify the dirstate parent.
                      '''
                      return self._parentwriters > 0
                  @propertycache
                  def _map(self):
                      '''Return the dirstate contents as a map from filename to
                      (state, mode, size, time).'''
                      self._read()
                      return self._map
                  @propertycache
                  def _copymap(self):
                      self._read()
                      return self._copymap
                  @propertycache
                  def _foldmap(self):
                      f = {}
                      normcase = util.normcase
                      for name, s in self._map.iteritems():
                          if s[0] != 'r':
                              f[normcase(name)] = name
                      for name in self._dirs:
                          f[normcase(name)] = name
                      f['.'] = '.' # prevents useless util.fspath() invocation
                      return f
                  @repocache('branch')
                  def _branch(self):
                      try:
                          return self._opener.read("branch").strip() or "default"
                      except IOError, inst:
                          if inst.errno != errno.ENOENT:
                              raise
                          return "default"
                  @propertycache
                  def _pl(self):
                      try:
                          fp = self._opener("dirstate")
                          st = fp.read(40)
                          fp.close()
                          l = len(st)
                          if l == 40:
                              return st[:20], st[20:40]
                          elif l > 0 and l < 40:
                              raise util.Abort(_('working directory state appears damaged!'))
                      except IOError, err:
                          if err.errno != errno.ENOENT:
                              raise
                      return [nullid, nullid]
                  @propertycache
                  def _dirs(self):
                      return scmutil.dirs(self._map, 'r')
                  def dirs(self):
                      return self._dirs
                  @rootcache('.hgignore')
                  def _ignore(self):
                      files = [self._join('.hgignore')]
                      for name, path in self._ui.configitems("ui"):
                          if name == 'ignore' or name.startswith('ignore.'):
                              # we need to use os.path.join here rather than self._join
                              # because path is arbitrary and user-specified
                              files.append(os.path.join(self._rootdir, util.expandpath(path)))
                      return ignore.ignore(self._root, files, self._ui.warn)
                  @propertycache
                  def _slash(self):
                      return self._ui.configbool('ui', 'slash') and os.sep != '/'
                  @propertycache
                  def _checklink(self):
                      return util.checklink(self._root)
                  @propertycache
                  def _checkexec(self):
                      return util.checkexec(self._root)
                  @propertycache
                  def _checkcase(self):
                      return not util.checkcase(self._join('.hg'))
                  def _join(self, f):
                      # much faster than os.path.join()
                      # it's safe because f is always a relative path
                      return self._rootdir + f
                  def flagfunc(self, buildfallback):
                      if self._checklink and self._checkexec:
                          def f(x):
                              try:
                                  st = os.lstat(self._join(x))
                                  if util.statislink(st):
                                      return 'l'
                                  if util.statisexec(st):
                                      return 'x'
                              except OSError:
                                  pass
                              return ''
                          return f
                      fallback = buildfallback()
                      if self._checklink:
                          def f(x):
                              if os.path.islink(self._join(x)):
                                  return 'l'
                              if 'x' in fallback(x):
                                  return 'x'
                              return ''
                          return f
                      if self._checkexec:
                          def f(x):
                              if 'l' in fallback(x):
                                  return 'l'
                              if util.isexec(self._join(x)):
                                  return 'x'
                              return ''
                          return f
                      else:
                          return fallback
                  @propertycache
                  def _cwd(self):
                      return os.getcwd()
                  def getcwd(self):
                      cwd = self._cwd
                      if cwd == self._root:
                          return ''
                      # self._root ends with a path separator if self._root is '/' or 'C:\'
                      rootsep = self._root
                      if not util.endswithsep(rootsep):
                          rootsep += os.sep
                      if cwd.startswith(rootsep):
                          return cwd[len(rootsep):]
                      else:
                          # we're outside the repo. return an absolute path.
                          return cwd
                  def pathto(self, f, cwd=None):
                      if cwd is None:
                          cwd = self.getcwd()
                      path = util.pathto(self._root, cwd, f)
                      if self._slash:
                          return util.pconvert(path)
                      return path
                  def __getitem__(self, key):
                      '''Return the current state of key (a filename) in the dirstate.
                      States are:
                        n  normal
                        m  needs merging
                        r  marked for removal
                        a  marked for addition
                        ?  not tracked
                      '''
                      return self._map.get(key, ("?",))[0]
                  def __contains__(self, key):
                      return key in self._map
                  def __iter__(self):
                      for x in sorted(self._map):
                          yield x
                  def iteritems(self):
                      return self._map.iteritems()
                  def parents(self):
                      return [self._validate(p) for p in self._pl]
                  def p1(self):
                      return self._validate(self._pl[0])
                  def p2(self):
                      return self._validate(self._pl[1])
                  def branch(self):
                      return encoding.tolocal(self._branch)
                  def setparents(self, p1, p2=nullid):
                      """Set dirstate parents to p1 and p2.
                      When moving from two parents to one, 'm' merged entries a
                      adjusted to normal and previous copy records discarded and
                      returned by the call.
                      See localrepo.setparents()
                      """
                      if self._parentwriters == 0:
                          raise ValueError("cannot set dirstate parent without "
                                           "calling dirstate.beginparentchange")
                      self._dirty = self._dirtypl = True
                      oldp2 = self._pl[1]
                      self._pl = p1, p2
                      copies = {}
                      if oldp2 != nullid and p2 == nullid:
                          for f, s in self._map.iteritems():
                              # Discard 'm' markers when moving away from a merge state
                              if s[0] == 'm':
                                  if f in self._copymap:
                                      copies[f] = self._copymap[f]
                                  self.normallookup(f)
                              # Also fix up otherparent markers
                              elif s[0] == 'n' and s[2] == -2:
                                  if f in self._copymap:
                                      copies[f] = self._copymap[f]
                                  self.add(f)
                      return copies
                  def setbranch(self, branch):
                      self._branch = encoding.fromlocal(branch)
                      f = self._opener('branch', 'w', atomictemp=True)
                      try:
                          f.write(self._branch + '\n')
                          f.close()
                          # make sure filecache has the correct stat info for _branch after
                          # replacing the underlying file
                          ce = self._filecache['_branch']
                          if ce:
                              ce.refresh()
                      except: # re-raises
                          f.discard()
                          raise
                  def _read(self):
                      self._map = {}
                      self._copymap = {}
                      try:
                          st = self._opener.read("dirstate")
                      except IOError, err:
                          if err.errno != errno.ENOENT:
                              raise
                          return
                      if not st:
                          return
                      # Python's garbage collector triggers a GC each time a certain number
                      # of container objects (the number being defined by
                      # gc.get_threshold()) are allocated. parse_dirstate creates a tuple
                      # for each file in the dirstate. The C version then immediately marks
                      # them as not to be tracked by the collector. However, this has no
                      # effect on when GCs are triggered, only on what objects the GC looks
                      # into. This means that O(number of files) GCs are unavoidable.
                      # Depending on when in the process's lifetime the dirstate is parsed,
                      # this can get very expensive. As a workaround, disable GC while
                      # parsing the dirstate.
                      #
                      # (we cannot decorate the function directly since it is in a C module)
                      parse_dirstate = util.nogc(parsers.parse_dirstate)
                      p = parse_dirstate(self._map, self._copymap, st)
                      if not self._dirtypl:
                          self._pl = p
                  def invalidate(self):
                      for a in ("_map", "_copymap", "_foldmap", "_branch", "_pl", "_dirs",
                              "_ignore"):
                          if a in self.__dict__:
                              delattr(self, a)
                      self._lastnormaltime = 0
                      self._dirty = False
                      self._parentwriters = 0
                  def copy(self, source, dest):
                      """Mark dest as a copy of source. Unmark dest if source is None."""
                      if source == dest:
                          return
                      self._dirty = True
                      if source is not None:
                          self._copymap[dest] = source
                      elif dest in self._copymap:
                          del self._copymap[dest]
                  def copied(self, file):
                      return self._copymap.get(file, None)
                  def copies(self):
                      return self._copymap
                  def _droppath(self, f):
                      if self[f] not in "?r" and "_dirs" in self.__dict__:
                          self._dirs.delpath(f)
                  def _addpath(self, f, state, mode, size, mtime):
                      oldstate = self[f]
                      if state == 'a' or oldstate == 'r':
                          scmutil.checkfilename(f)
                          if f in self._dirs:
                              raise util.Abort(_('directory %r already in dirstate') % f)
                          # shadows
                          for d in scmutil.finddirs(f):
                              if d in self._dirs:
                                  break
                              if d in self._map and self[d] != 'r':
                                  raise util.Abort(
                                      _('file %r in dirstate clashes with %r') % (d, f))
                      if oldstate in "?r" and "_dirs" in self.__dict__:
                          self._dirs.addpath(f)
                      self._dirty = True
                      self._map[f] = dirstatetuple(state, mode, size, mtime)
                  def normal(self, f):
                      '''Mark a file normal and clean.'''
                      s = os.lstat(self._join(f))
                      mtime = int(s.st_mtime)
                      self._addpath(f, 'n', s.st_mode,
                                    s.st_size & _rangemask, mtime & _rangemask)
                      if f in self._copymap:
                          del self._copymap[f]
                      if mtime > self._lastnormaltime:
                          # Remember the most recent modification timeslot for status(),
                          # to make sure we won't miss future size-preserving file content
                          # modifications that happen within the same timeslot.
                          self._lastnormaltime = mtime
                  def normallookup(self, f):
                      '''Mark a file normal, but possibly dirty.'''
                      if self._pl[1] != nullid and f in self._map:
                          # if there is a merge going on and the file was either
                          # in state 'm' (-1) or coming from other parent (-2) before
                          # being removed, restore that state.
                          entry = self._map[f]
                          if entry[0] == 'r' and entry[2] in (-1, -2):
                              source = self._copymap.get(f)
                              if entry[2] == -1:
                                  self.merge(f)
                              elif entry[2] == -2:
                                  self.otherparent(f)
                              if source:
                                  self.copy(source, f)
                              return
                          if entry[0] == 'm' or entry[0] == 'n' and entry[2] == -2:
                              return
                      self._addpath(f, 'n', 0, -1, -1)
                      if f in self._copymap:
                          del self._copymap[f]
                  def otherparent(self, f):
                      '''Mark as coming from the other parent, always dirty.'''
                      if self._pl[1] == nullid:
                          raise util.Abort(_("setting %r to other parent "
                                             "only allowed in merges") % f)
                      if f in self and self[f] == 'n':
                          # merge-like
                          self._addpath(f, 'm', 0, -2, -1)
                      else:
                          # add-like
                          self._addpath(f, 'n', 0, -2, -1)
                      if f in self._copymap:
                          del self._copymap[f]
                  def add(self, f):
                      '''Mark a file added.'''
                      self._addpath(f, 'a', 0, -1, -1)
                      if f in self._copymap:
                          del self._copymap[f]
                  def remove(self, f):
                      '''Mark a file removed.'''
                      self._dirty = True
                      self._droppath(f)
                      size = 0
                      if self._pl[1] != nullid and f in self._map:
                          # backup the previous state
                          entry = self._map[f]
                          if entry[0] == 'm': # merge
                              size = -1
                          elif entry[0] == 'n' and entry[2] == -2: # other parent
                              size = -2
                      self._map[f] = dirstatetuple('r', 0, size, 0)
                      if size == 0 and f in self._copymap:
                          del self._copymap[f]
                  def merge(self, f):
                      '''Mark a file merged.'''
                      if self._pl[1] == nullid:
                          return self.normallookup(f)
                      return self.otherparent(f)
                  def drop(self, f):
                      '''Drop a file from the dirstate'''
                      if f in self._map:
                          self._dirty = True
                          self._droppath(f)
                          del self._map[f]
                  def _normalize(self, path, isknown, ignoremissing=False, exists=None):
                      normed = util.normcase(path)
                      folded = self._foldmap.get(normed, None)
                      if folded is None:
                          if isknown:
                              folded = path
                          else:
                              if exists is None:
                                  exists = os.path.lexists(os.path.join(self._root, path))
                              if not exists:
                                  # Maybe a path component exists
                                  if not ignoremissing and '/' in path:
                                      d, f = path.rsplit('/', 1)
                                      d = self._normalize(d, isknown, ignoremissing, None)
                                      folded = d + "/" + f
                                  else:
                                      # No path components, preserve original case
                                      folded = path
                              else:
                                  # recursively normalize leading directory components
                                  # against dirstate
                                  if '/' in normed:
                                      d, f = normed.rsplit('/', 1)
                                      d = self._normalize(d, isknown, ignoremissing, True)
                                      r = self._root + "/" + d
                                      folded = d + "/" + util.fspath(f, r)
                                  else:
                                      folded = util.fspath(normed, self._root)
                                  self._foldmap[normed] = folded
                      return folded
                  def normalize(self, path, isknown=False, ignoremissing=False):
                      '''
                      normalize the case of a pathname when on a casefolding filesystem
                      isknown specifies whether the filename came from walking the
                      disk, to avoid extra filesystem access.
                      If ignoremissing is True, missing path are returned
                      unchanged. Otherwise, we try harder to normalize possibly
                      existing path components.
                      The normalized case is determined based on the following precedence:
                      - version of name already stored in the dirstate
                      - version of name stored on disk
                      - version provided via command arguments
                      '''
                      if self._checkcase:
                          return self._normalize(path, isknown, ignoremissing)
                      return path
                  def clear(self):
                      self._map = {}
                      if "_dirs" in self.__dict__:
                          delattr(self, "_dirs")
                      self._copymap = {}
                      self._pl = [nullid, nullid]
                      self._lastnormaltime = 0
                      self._dirty = True
                  def rebuild(self, parent, allfiles, changedfiles=None):
                      changedfiles = changedfiles or allfiles
                      oldmap = self._map
                      self.clear()
                      for f in allfiles:
                          if f not in changedfiles:
                              self._map[f] = oldmap[f]
                          else:
                              if 'x' in allfiles.flags(f):
                                  self._map[f] = dirstatetuple('n', 0777, -1, 0)
                              else:
                                  self._map[f] = dirstatetuple('n', 0666, -1, 0)
                      self._pl = (parent, nullid)
                      self._dirty = True
                  def write(self):
                      if not self._dirty:
                          return
                      # enough 'delaywrite' prevents 'pack_dirstate' from dropping
                      # timestamp of each entries in dirstate, because of 'now > mtime'
                      delaywrite = self._ui.configint('debug', 'dirstate.delaywrite', 0)
                      if delaywrite > 0:
                          import time # to avoid useless import
                          time.sleep(delaywrite)
                      st = self._opener("dirstate", "w", atomictemp=True)
                      # use the modification time of the newly created temporary file as the
                      # filesystem's notion of 'now'
                      now = util.fstat(st).st_mtime
                      st.write(parsers.pack_dirstate(self._map, self._copymap, self._pl, now))
                      st.close()
                      self._lastnormaltime = 0
                      self._dirty = self._dirtypl = False
                  def _dirignore(self, f):
                      if f == '.':
                          return False
                      if self._ignore(f):
                          return True
                      for p in scmutil.finddirs(f):
                          if self._ignore(p):
                              return True
                      return False
                  def _walkexplicit(self, match, subrepos):
                      '''Get stat data about the files explicitly specified by match.
                      Return a triple (results, dirsfound, dirsnotfound).
                      - results is a mapping from filename to stat result. It also contains
                        listings mapping subrepos and .hg to None.
                      - dirsfound is a list of files found to be directories.
                      - dirsnotfound is a list of files that the dirstate thinks are
                        directories and that were not found.'''
                      def badtype(mode):
                          kind = _('unknown')
                          if stat.S_ISCHR(mode):
                              kind = _('character device')
                          elif stat.S_ISBLK(mode):
                              kind = _('block device')
                          elif stat.S_ISFIFO(mode):
                              kind = _('fifo')
                          elif stat.S_ISSOCK(mode):
                              kind = _('socket')
                          elif stat.S_ISDIR(mode):
                              kind = _('directory')
                          return _('unsupported file type (type is %s)') % kind
                      matchedir = match.explicitdir
                      badfn = match.bad
                      dmap = self._map
                      normpath = util.normpath
                      lstat = os.lstat
                      getkind = stat.S_IFMT
                      dirkind = stat.S_IFDIR
                      regkind = stat.S_IFREG
                      lnkkind = stat.S_IFLNK
                      join = self._join
                      dirsfound = []
                      foundadd = dirsfound.append
                      dirsnotfound = []
                      notfoundadd = dirsnotfound.append
-                     if match.matchfn != match.exact and self._checkcase:
+                     if not match.isexact() and self._checkcase:
                          normalize = self._normalize
                      else:
                          normalize = None
                      files = sorted(match.files())
                      subrepos.sort()
                      i, j = 0, 0
                      while i < len(files) and j < len(subrepos):
                          subpath = subrepos[j] + "/"
                          if files[i] < subpath:
                              i += 1
                              continue
                          while i < len(files) and files[i].startswith(subpath):
                              del files[i]
                          j += 1
                      if not files or '.' in files:
                          files = ['']
                      results = dict.fromkeys(subrepos)
                      results['.hg'] = None
                      alldirs = None
                      for ff in files:
                          if normalize:
                              nf = normalize(normpath(ff), False, True)
                          else:
                              nf = normpath(ff)
                          if nf in results:
                              continue
                          try:
                              st = lstat(join(nf))
                              kind = getkind(st.st_mode)
                              if kind == dirkind:
                                  if nf in dmap:
                                      # file replaced by dir on disk but still in dirstate
                                      results[nf] = None
                                  if matchedir:
                                      matchedir(nf)
                                  foundadd(nf)
                              elif kind == regkind or kind == lnkkind:
                                  results[nf] = st
                              else:
                                  badfn(ff, badtype(kind))
                                  if nf in dmap:
                                      results[nf] = None
                          except OSError, inst: # nf not found on disk - it is dirstate only
                              if nf in dmap: # does it exactly match a missing file?
                                  results[nf] = None
                              else: # does it match a missing directory?
                                  if alldirs is None:
                                      alldirs = scmutil.dirs(dmap)
                                  if nf in alldirs:
                                      if matchedir:
                                          matchedir(nf)
                                      notfoundadd(nf)
                                  else:
                                      badfn(ff, inst.strerror)
                      return results, dirsfound, dirsnotfound
                  def walk(self, match, subrepos, unknown, ignored, full=True):
                      '''
                      Walk recursively through the directory tree, finding all files
                      matched by match.
                      If full is False, maybe skip some known-clean files.
                      Return a dict mapping filename to stat-like object (either
                      mercurial.osutil.stat instance or return value of os.stat()).
                      '''
                      # full is a flag that extensions that hook into walk can use -- this
                      # implementation doesn't use it at all. This satisfies the contract
                      # because we only guarantee a "maybe".
                      if ignored:
                          ignore = util.never
                          dirignore = util.never
                      elif unknown:
                          ignore = self._ignore
                          dirignore = self._dirignore
                      else:
                          # if not unknown and not ignored, drop dir recursion and step 2
                          ignore = util.always
                          dirignore = util.always
                      matchfn = match.matchfn
                      matchalways = match.always()
                      matchtdir = match.traversedir
                      dmap = self._map
                      listdir = osutil.listdir
                      lstat = os.lstat
                      dirkind = stat.S_IFDIR
                      regkind = stat.S_IFREG
                      lnkkind = stat.S_IFLNK
                      join = self._join
                      exact = skipstep3 = False
-                     if matchfn == match.exact: # match.exact
+                     if match.isexact(): # match.exact
                          exact = True
                          dirignore = util.always # skip step 2
                      elif match.files() and not match.anypats(): # match.match, no patterns
                          skipstep3 = True
                      if not exact and self._checkcase:
                          normalize = self._normalize
                          skipstep3 = False
                      else:
                          normalize = None
                      # step 1: find all explicit files
                      results, work, dirsnotfound = self._walkexplicit(match, subrepos)
                      skipstep3 = skipstep3 and not (work or dirsnotfound)
                      work = [d for d in work if not dirignore(d)]
                      wadd = work.append
                      # step 2: visit subdirectories
                      while work:
                          nd = work.pop()
                          skip = None
                          if nd == '.':
                              nd = ''
                          else:
                              skip = '.hg'
                          try:
                              entries = listdir(join(nd), stat=True, skip=skip)
                          except OSError, inst:
                              if inst.errno in (errno.EACCES, errno.ENOENT):
                                  match.bad(self.pathto(nd), inst.strerror)
                                  continue
                              raise
                          for f, kind, st in entries:
                              if normalize:
                                  nf = normalize(nd and (nd + "/" + f) or f, True, True)
                              else:
                                  nf = nd and (nd + "/" + f) or f
                              if nf not in results:
                                  if kind == dirkind:
                                      if not ignore(nf):
                                          if matchtdir:
                                              matchtdir(nf)
                                          wadd(nf)
                                      if nf in dmap and (matchalways or matchfn(nf)):
                                          results[nf] = None
                                  elif kind == regkind or kind == lnkkind:
                                      if nf in dmap:
                                          if matchalways or matchfn(nf):
                                              results[nf] = st
                                      elif (matchalways or matchfn(nf)) and not ignore(nf):
                                          results[nf] = st
                                  elif nf in dmap and (matchalways or matchfn(nf)):
                                      results[nf] = None
                      for s in subrepos:
                          del results[s]
                      del results['.hg']
                      # step 3: visit remaining files from dmap
                      if not skipstep3 and not exact:
                          # If a dmap file is not in results yet, it was either
                          # a) not matching matchfn b) ignored, c) missing, or d) under a
                          # symlink directory.
                          if not results and matchalways:
                              visit = dmap.keys()
                          else:
                              visit = [f for f in dmap if f not in results and matchfn(f)]
                          visit.sort()
                          if unknown:
                              # unknown == True means we walked all dirs under the roots
                              # that wasn't ignored, and everything that matched was stat'ed
                              # and is already in results.
                              # The rest must thus be ignored or under a symlink.
                              audit_path = pathutil.pathauditor(self._root)
                              for nf in iter(visit):
                                  # Report ignored items in the dmap as long as they are not
                                  # under a symlink directory.
                                  if audit_path.check(nf):
                                      try:
                                          results[nf] = lstat(join(nf))
                                          # file was just ignored, no links, and exists
                                      except OSError:
                                          # file doesn't exist
                                          results[nf] = None
                                  else:
                                      # It's either missing or under a symlink directory
                                      # which we in this case report as missing
                                      results[nf] = None
                          else:
                              # We may not have walked the full directory tree above,
                              # so stat and check everything we missed.
                              nf = iter(visit).next
                              for st in util.statfiles([join(i) for i in visit]):
                                  results[nf()] = st
                      return results
                  def status(self, match, subrepos, ignored, clean, unknown):
                      '''Determine the status of the working copy relative to the
                      dirstate and return a pair of (unsure, status), where status is of type
                      scmutil.status and:
                        unsure:
                          files that might have been modified since the dirstate was
                          written, but need to be read to be sure (size is the same
                          but mtime differs)
                        status.modified:
                          files that have definitely been modified since the dirstate
                          was written (different size or mode)
                        status.clean:
                          files that have definitely not been modified since the
                          dirstate was written
                      '''
                      listignored, listclean, listunknown = ignored, clean, unknown
                      lookup, modified, added, unknown, ignored = [], [], [], [], []
                      removed, deleted, clean = [], [], []
                      dmap = self._map
                      ladd = lookup.append            # aka "unsure"
                      madd = modified.append
                      aadd = added.append
                      uadd = unknown.append
                      iadd = ignored.append
                      radd = removed.append
                      dadd = deleted.append
                      cadd = clean.append
                      mexact = match.exact
                      dirignore = self._dirignore
                      checkexec = self._checkexec
                      copymap = self._copymap
                      lastnormaltime = self._lastnormaltime
                      # We need to do full walks when either
                      # - we're listing all clean files, or
                      # - match.traversedir does something, because match.traversedir should
                      #   be called for every dir in the working dir
                      full = listclean or match.traversedir is not None
                      for fn, st in self.walk(match, subrepos, listunknown, listignored,
                                              full=full).iteritems():
                          if fn not in dmap:
                              if (listignored or mexact(fn)) and dirignore(fn):
                                  if listignored:
                                      iadd(fn)
                              else:
                                  uadd(fn)
                              continue
                          # This is equivalent to 'state, mode, size, time = dmap[fn]' but not
                          # written like that for performance reasons. dmap[fn] is not a
                          # Python tuple in compiled builds. The CPython UNPACK_SEQUENCE
                          # opcode has fast paths when the value to be unpacked is a tuple or
                          # a list, but falls back to creating a full-fledged iterator in
                          # general. That is much slower than simply accessing and storing the
                          # tuple members one by one.
                          t = dmap[fn]
                          state = t[0]
                          mode = t[1]
                          size = t[2]
                          time = t[3]
                          if not st and state in "nma":
                              dadd(fn)
                          elif state == 'n':
                              mtime = int(st.st_mtime)
                              if (size >= 0 and
                                  ((size != st.st_size and size != st.st_size & _rangemask)
                                   or ((mode ^ st.st_mode) & 0100 and checkexec))
                                  or size == -2 # other parent
                                  or fn in copymap):
                                  madd(fn)
                              elif time != mtime and time != mtime & _rangemask:
                                  ladd(fn)
                              elif mtime == lastnormaltime:
                                  # fn may have just been marked as normal and it may have
                                  # changed in the same second without changing its size.
                                  # This can happen if we quickly do multiple commits.
                                  # Force lookup, so we don't miss such a racy file change.
                                  ladd(fn)
                              elif listclean:
                                  cadd(fn)
                          elif state == 'm':
                              madd(fn)
                          elif state == 'a':
                              aadd(fn)
                          elif state == 'r':
                              radd(fn)
                      return (lookup, scmutil.status(modified, added, removed, deleted,
                                                     unknown, ignored, clean))
                  def matches(self, match):
                      '''
                      return files in the dirstate (in whatever state) filtered by match
                      '''
                      dmap = self._map
                      if match.always():
                          return dmap.keys()
                      files = match.files()
-                     if match.matchfn == match.exact:
+                     if match.isexact():
                          # fast path -- filter the other way around, since typically files is
                          # much smaller than dmap
                          return [f for f in files if f in dmap]
                      if not match.anypats() and util.all(fn in dmap for fn in files):
                          # fast path -- all the values are known to be files, so just return
                          # that
                          return list(files)
                      return [f for f in dmap if match(f)]

mercurial/manifest.py

0 +2 -2

              # manifest.py - manifest revision class for mercurial
              #
              # Copyright 2005-2007 Matt Mackall <mpm@selenic.com>
              #
              # This software may be used and distributed according to the terms of the
              # GNU General Public License version 2 or any later version.
              from i18n import _
              import mdiff, parsers, error, revlog, util, scmutil
              import array, struct
              propertycache = util.propertycache
              class _lazymanifest(dict):
                  """This is the pure implementation of lazymanifest.
                  It has not been optimized *at all* and is not lazy.
                  """
                  def __init__(self, data):
                      # This init method does a little bit of excessive-looking
                      # precondition checking. This is so that the behavior of this
                      # class exactly matches its C counterpart to try and help
                      # prevent surprise breakage for anyone that develops against
                      # the pure version.
                      if data and data[-1] != '\n':
                          raise ValueError('Manifest did not end in a newline.')
                      dict.__init__(self)
                      prev = None
                      for l in data.splitlines():
                          if prev is not None and prev > l:
                              raise ValueError('Manifest lines not in sorted order.')
                          prev = l
                          f, n = l.split('\0')
                          if len(n) > 40:
                              self[f] = revlog.bin(n[:40]), n[40:]
                          else:
                              self[f] = revlog.bin(n), ''
                  def __setitem__(self, k, v):
                      node, flag = v
                      assert node is not None
                      if len(node) > 21:
                          node = node[:21] # match c implementation behavior
                      dict.__setitem__(self, k, (node, flag))
                  def __iter__(self):
                      return iter(sorted(dict.keys(self)))
                  def iterkeys(self):
                      return iter(sorted(dict.keys(self)))
                  def iterentries(self):
                      return ((f, e[0], e[1]) for f, e in sorted(self.iteritems()))
                  def copy(self):
                      c = _lazymanifest('')
                      c.update(self)
                      return c
                  def diff(self, m2, clean=False):
                      '''Finds changes between the current manifest and m2.'''
                      diff = {}
                      for fn, e1 in self.iteritems():
                          if fn not in m2:
                              diff[fn] = e1, (None, '')
                          else:
                              e2 = m2[fn]
                              if e1 != e2:
                                  diff[fn] = e1, e2
                              elif clean:
                                  diff[fn] = None
                      for fn, e2 in m2.iteritems():
                          if fn not in self:
                              diff[fn] = (None, ''), e2
                      return diff
                  def filtercopy(self, filterfn):
                      c = _lazymanifest('')
                      for f, n, fl in self.iterentries():
                          if filterfn(f):
                              c[f] = n, fl
                      return c
                  def text(self):
                      """Get the full data of this manifest as a bytestring."""
                      fl = sorted(self.iterentries())
                      _hex = revlog.hex
                      # if this is changed to support newlines in filenames,
                      # be sure to check the templates/ dir again (especially *-raw.tmpl)
                      return ''.join("%s\0%s%s\n" % (
                          f, _hex(n[:20]), flag) for f, n, flag in fl)
              try:
                  _lazymanifest = parsers.lazymanifest
              except AttributeError:
                  pass
              class manifestdict(object):
                  def __init__(self, data=''):
                      self._lm = _lazymanifest(data)
                  def __getitem__(self, key):
                      return self._lm[key][0]
                  def find(self, key):
                      return self._lm[key]
                  def __len__(self):
                      return len(self._lm)
                  def __setitem__(self, key, node):
                      self._lm[key] = node, self.flags(key, '')
                  def __contains__(self, key):
                      return key in self._lm
                  def __delitem__(self, key):
                      del self._lm[key]
                  def __iter__(self):
                      return self._lm.__iter__()
                  def iterkeys(self):
                      return self._lm.iterkeys()
                  def keys(self):
                      return list(self.iterkeys())
                  def intersectfiles(self, files):
                      '''make a new lazymanifest with the intersection of self with files
                      The algorithm assumes that files is much smaller than self.'''
                      ret = manifestdict()
                      lm = self._lm
                      for fn in files:
                          if fn in lm:
                              ret._lm[fn] = self._lm[fn]
                      return ret
                  def filesnotin(self, m2):
                      '''Set of files in this manifest that are not in the other'''
                      files = set(self)
                      files.difference_update(m2)
                      return files
                  @propertycache
                  def _dirs(self):
                      return scmutil.dirs(self)
                  def dirs(self):
                      return self._dirs
                  def hasdir(self, dir):
                      return dir in self._dirs
                  def matches(self, match):
                      '''generate a new manifest filtered by the match argument'''
                      if match.always():
                          return self.copy()
                      files = match.files()
-                     if (len(files) < 100 and (match.matchfn == match.exact or
+                     if (len(files) < 100 and (match.isexact() or
                          (not match.anypats() and util.all(fn in self for fn in files)))):
                          return self.intersectfiles(files)
                      lm = manifestdict('')
                      lm._lm = self._lm.filtercopy(match)
                      return lm
                  def diff(self, m2, clean=False):
                      '''Finds changes between the current manifest and m2.
                      Args:
                        m2: the manifest to which this manifest should be compared.
                        clean: if true, include files unchanged between these manifests
                               with a None value in the returned dictionary.
                      The result is returned as a dict with filename as key and
                      values of the form ((n1,fl1),(n2,fl2)), where n1/n2 is the
                      nodeid in the current/other manifest and fl1/fl2 is the flag
                      in the current/other manifest. Where the file does not exist,
                      the nodeid will be None and the flags will be the empty
                      string.
                      '''
                      return self._lm.diff(m2._lm, clean)
                  def setflag(self, key, flag):
                      self._lm[key] = self[key], flag
                  def get(self, key, default=None):
                      try:
                          return self._lm[key][0]
                      except KeyError:
                          return default
                  def flags(self, key, default=''):
                      try:
                          return self._lm[key][1]
                      except KeyError:
                          return default
                  def copy(self):
                      c = manifestdict('')
                      c._lm = self._lm.copy()
                      return c
                  def iteritems(self):
                      return (x[:2] for x in self._lm.iterentries())
                  def text(self):
                      return self._lm.text()
                  def fastdelta(self, base, changes):
                      """Given a base manifest text as an array.array and a list of changes
                      relative to that text, compute a delta that can be used by revlog.
                      """
                      delta = []
                      dstart = None
                      dend = None
                      dline = [""]
                      start = 0
                      # zero copy representation of base as a buffer
                      addbuf = util.buffer(base)
                      # start with a readonly loop that finds the offset of
                      # each line and creates the deltas
                      for f, todelete in changes:
                          # bs will either be the index of the item or the insert point
                          start, end = _msearch(addbuf, f, start)
                          if not todelete:
                              h, fl = self._lm[f]
                              l = "%s\0%s%s\n" % (f, revlog.hex(h), fl)
                          else:
                              if start == end:
                                  # item we want to delete was not found, error out
                                  raise AssertionError(
                                          _("failed to remove %s from manifest") % f)
                              l = ""
                          if dstart is not None and dstart <= start and dend >= start:
                              if dend < end:
                                  dend = end
                              if l:
                                  dline.append(l)
                          else:
                              if dstart is not None:
                                  delta.append([dstart, dend, "".join(dline)])
                              dstart = start
                              dend = end
                              dline = [l]
                      if dstart is not None:
                          delta.append([dstart, dend, "".join(dline)])
                      # apply the delta to the base, and get a delta for addrevision
                      deltatext, arraytext = _addlistdelta(base, delta)
                      return arraytext, deltatext
              def _msearch(m, s, lo=0, hi=None):
                  '''return a tuple (start, end) that says where to find s within m.
                  If the string is found m[start:end] are the line containing
                  that string.  If start == end the string was not found and
                  they indicate the proper sorted insertion point.
                  m should be a buffer or a string
                  s is a string'''
                  def advance(i, c):
                      while i < lenm and m[i] != c:
                          i += 1
                      return i
                  if not s:
                      return (lo, lo)
                  lenm = len(m)
                  if not hi:
                      hi = lenm
                  while lo < hi:
                      mid = (lo + hi) // 2
                      start = mid
                      while start > 0 and m[start - 1] != '\n':
                          start -= 1
                      end = advance(start, '\0')
                      if m[start:end] < s:
                          # we know that after the null there are 40 bytes of sha1
                          # this translates to the bisect lo = mid + 1
                          lo = advance(end + 40, '\n') + 1
                      else:
                          # this translates to the bisect hi = mid
                          hi = start
                  end = advance(lo, '\0')
                  found = m[lo:end]
                  if s == found:
                      # we know that after the null there are 40 bytes of sha1
                      end = advance(end + 40, '\n')
                      return (lo, end + 1)
                  else:
                      return (lo, lo)
              def _checkforbidden(l):
                  """Check filenames for illegal characters."""
                  for f in l:
                      if '\n' in f or '\r' in f:
                          raise error.RevlogError(
                              _("'\\n' and '\\r' disallowed in filenames: %r") % f)
              # apply the changes collected during the bisect loop to our addlist
              # return a delta suitable for addrevision
              def _addlistdelta(addlist, x):
                  # for large addlist arrays, building a new array is cheaper
                  # than repeatedly modifying the existing one
                  currentposition = 0
                  newaddlist = array.array('c')
                  for start, end, content in x:
                      newaddlist += addlist[currentposition:start]
                      if content:
                          newaddlist += array.array('c', content)
                      currentposition = end
                  newaddlist += addlist[currentposition:]
                  deltatext = "".join(struct.pack(">lll", start, end, len(content))
                                 + content for start, end, content in x)
                  return deltatext, newaddlist
              def _splittopdir(f):
                  if '/' in f:
                      dir, subpath = f.split('/', 1)
                      return dir + '/', subpath
                  else:
                      return '', f
              class treemanifest(object):
                  def __init__(self, dir='', text=''):
                      self._dir = dir
                      self._dirs = {}
                      # Using _lazymanifest here is a little slower than plain old dicts
                      self._files = {}
                      self._flags = {}
                      lm = _lazymanifest(text)
                      for f, n, fl in lm.iterentries():
                          self[f] = n
                          if fl:
                              self.setflag(f, fl)
                  def _subpath(self, path):
                      return self._dir + path
                  def __len__(self):
                      size = len(self._files)
                      for m in self._dirs.values():
                          size += m.__len__()
                      return size
                  def __str__(self):
                      return '<treemanifest dir=%s>' % self._dir
                  def iteritems(self):
                      for p, n in sorted(self._dirs.items() + self._files.items()):
                          if p in self._files:
                              yield self._subpath(p), n
                          else:
                              for f, sn in n.iteritems():
                                  yield f, sn
                  def iterkeys(self):
                      for p in sorted(self._dirs.keys() + self._files.keys()):
                          if p in self._files:
                              yield self._subpath(p)
                          else:
                              for f in self._dirs[p].iterkeys():
                                  yield f
                  def keys(self):
                      return list(self.iterkeys())
                  def __iter__(self):
                      return self.iterkeys()
                  def __contains__(self, f):
                      if f is None:
                          return False
                      dir, subpath = _splittopdir(f)
                      if dir:
                          if dir not in self._dirs:
                              return False
                          return self._dirs[dir].__contains__(subpath)
                      else:
                          return f in self._files
                  def get(self, f, default=None):
                      dir, subpath = _splittopdir(f)
                      if dir:
                          if dir not in self._dirs:
                              return default
                          return self._dirs[dir].get(subpath, default)
                      else:
                          return self._files.get(f, default)
                  def __getitem__(self, f):
                      dir, subpath = _splittopdir(f)
                      if dir:
                          return self._dirs[dir].__getitem__(subpath)
                      else:
                          return self._files[f]
                  def flags(self, f):
                      dir, subpath = _splittopdir(f)
                      if dir:
                          if dir not in self._dirs:
                              return ''
                          return self._dirs[dir].flags(subpath)
                      else:
                          if f in self._dirs:
                              return ''
                          return self._flags.get(f, '')
                  def find(self, f):
                      dir, subpath = _splittopdir(f)
                      if dir:
                          return self._dirs[dir].find(subpath)
                      else:
                          return self._files[f], self._flags.get(f, '')
                  def __delitem__(self, f):
                      dir, subpath = _splittopdir(f)
                      if dir:
                          self._dirs[dir].__delitem__(subpath)
                          # If the directory is now empty, remove it
                          if not self._dirs[dir]._dirs and not self._dirs[dir]._files:
                              del self._dirs[dir]
                      else:
                          del self._files[f]
                          if f in self._flags:
                              del self._flags[f]
                  def __setitem__(self, f, n):
                      assert n is not None
                      dir, subpath = _splittopdir(f)
                      if dir:
                          if dir not in self._dirs:
                              self._dirs[dir] = treemanifest(self._subpath(dir))
                          self._dirs[dir].__setitem__(subpath, n)
                      else:
                          self._files[f] = n
                  def setflag(self, f, flags):
                      """Set the flags (symlink, executable) for path f."""
                      dir, subpath = _splittopdir(f)
                      if dir:
                          if dir not in self._dirs:
                              self._dirs[dir] = treemanifest(self._subpath(dir))
                          self._dirs[dir].setflag(subpath, flags)
                      else:
                          self._flags[f] = flags
                  def copy(self):
                      copy = treemanifest(self._dir)
                      for d in self._dirs:
                          copy._dirs[d] = self._dirs[d].copy()
                      copy._files = dict.copy(self._files)
                      copy._flags = dict.copy(self._flags)
                      return copy
                  def intersectfiles(self, files):
                      '''make a new treemanifest with the intersection of self with files
                      The algorithm assumes that files is much smaller than self.'''
                      ret = treemanifest()
                      for fn in files:
                          if fn in self:
                              ret[fn] = self[fn]
                              flags = self.flags(fn)
                              if flags:
                                  ret.setflag(fn, flags)
                      return ret
                  def filesnotin(self, m2):
                      '''Set of files in this manifest that are not in the other'''
                      files = set()
                      def _filesnotin(t1, t2):
                          for d, m1 in t1._dirs.iteritems():
                              if d in t2._dirs:
                                  m2 = t2._dirs[d]
                                  _filesnotin(m1, m2)
                              else:
                                  files.update(m1.iterkeys())
                          for fn in t1._files.iterkeys():
                              if fn not in t2._files:
                                  files.add(t1._subpath(fn))
                      _filesnotin(self, m2)
                      return files
                  @propertycache
                  def _alldirs(self):
                      return scmutil.dirs(self)
                  def dirs(self):
                      return self._alldirs
                  def hasdir(self, dir):
                      topdir, subdir = _splittopdir(dir)
                      if topdir:
                          if topdir in self._dirs:
                              return self._dirs[topdir].hasdir(subdir)
                          return False
                      return (dir + '/') in self._dirs
                  def matches(self, match):
                      '''generate a new manifest filtered by the match argument'''
                      if match.always():
                          return self.copy()
                      files = match.files()
-                     if (match.matchfn == match.exact or
+                     if (match.isexact() or
                          (not match.anypats() and util.all(fn in self for fn in files))):
                          return self.intersectfiles(files)
                      m = self.copy()
                      for fn in m.keys():
                          if not match(fn):
                              del m[fn]
                      return m
                  def diff(self, m2, clean=False):
                      '''Finds changes between the current manifest and m2.
                      Args:
                        m2: the manifest to which this manifest should be compared.
                        clean: if true, include files unchanged between these manifests
                               with a None value in the returned dictionary.
                      The result is returned as a dict with filename as key and
                      values of the form ((n1,fl1),(n2,fl2)), where n1/n2 is the
                      nodeid in the current/other manifest and fl1/fl2 is the flag
                      in the current/other manifest. Where the file does not exist,
                      the nodeid will be None and the flags will be the empty
                      string.
                      '''
                      result = {}
                      emptytree = treemanifest()
                      def _diff(t1, t2):
                          for d, m1 in t1._dirs.iteritems():
                              m2 = t2._dirs.get(d, emptytree)
                              _diff(m1, m2)
                          for d, m2 in t2._dirs.iteritems():
                              if d not in t1._dirs:
                                  _diff(emptytree, m2)
                          for fn, n1 in t1._files.iteritems():
                              fl1 = t1._flags.get(fn, '')
                              n2 = t2._files.get(fn, None)
                              fl2 = t2._flags.get(fn, '')
                              if n1 != n2 or fl1 != fl2:
                                  result[t1._subpath(fn)] = ((n1, fl1), (n2, fl2))
                              elif clean:
                                  result[t1._subpath(fn)] = None
                          for fn, n2 in t2._files.iteritems():
                              if fn not in t1._files:
                                  fl2 = t2._flags.get(fn, '')
                                  result[t2._subpath(fn)] = ((None, ''), (n2, fl2))
                      _diff(self, m2)
                      return result
                  def text(self):
                      """Get the full data of this manifest as a bytestring."""
                      fl = self.keys()
                      _checkforbidden(fl)
                      hex, flags = revlog.hex, self.flags
                      # if this is changed to support newlines in filenames,
                      # be sure to check the templates/ dir again (especially *-raw.tmpl)
                      return ''.join("%s\0%s%s\n" % (f, hex(self[f]), flags(f)) for f in fl)
              class manifest(revlog.revlog):
                  def __init__(self, opener):
                      # During normal operations, we expect to deal with not more than four
                      # revs at a time (such as during commit --amend). When rebasing large
                      # stacks of commits, the number can go up, hence the config knob below.
                      cachesize = 4
                      usetreemanifest = False
                      opts = getattr(opener, 'options', None)
                      if opts is not None:
                          cachesize = opts.get('manifestcachesize', cachesize)
                          usetreemanifest = opts.get('usetreemanifest', usetreemanifest)
                      self._mancache = util.lrucachedict(cachesize)
                      revlog.revlog.__init__(self, opener, "00manifest.i")
                      self._usetreemanifest = usetreemanifest
                  def _newmanifest(self, data=''):
                      if self._usetreemanifest:
                          return treemanifest('', data)
                      return manifestdict(data)
                  def readdelta(self, node):
                      r = self.rev(node)
                      d = mdiff.patchtext(self.revdiff(self.deltaparent(r), r))
                      return self._newmanifest(d)
                  def readfast(self, node):
                      '''use the faster of readdelta or read'''
                      r = self.rev(node)
                      deltaparent = self.deltaparent(r)
                      if deltaparent != revlog.nullrev and deltaparent in self.parentrevs(r):
                          return self.readdelta(node)
                      return self.read(node)
                  def read(self, node):
                      if node == revlog.nullid:
                          return self._newmanifest() # don't upset local cache
                      if node in self._mancache:
                          return self._mancache[node][0]
                      text = self.revision(node)
                      arraytext = array.array('c', text)
                      m = self._newmanifest(text)
                      self._mancache[node] = (m, arraytext)
                      return m
                  def find(self, node, f):
                      '''look up entry for a single file efficiently.
                      return (node, flags) pair if found, (None, None) if not.'''
                      m = self.read(node)
                      try:
                          return m.find(f)
                      except KeyError:
                          return None, None
                  def add(self, m, transaction, link, p1, p2, added, removed):
                      if p1 in self._mancache and not self._usetreemanifest:
                          # If our first parent is in the manifest cache, we can
                          # compute a delta here using properties we know about the
                          # manifest up-front, which may save time later for the
                          # revlog layer.
                          _checkforbidden(added)
                          # combine the changed lists into one list for sorting
                          work = [(x, False) for x in added]
                          work.extend((x, True) for x in removed)
                          # this could use heapq.merge() (from Python 2.6+) or equivalent
                          # since the lists are already sorted
                          work.sort()
                          arraytext, deltatext = m.fastdelta(self._mancache[p1][1], work)
                          cachedelta = self.rev(p1), deltatext
                          text = util.buffer(arraytext)
                      else:
                          # The first parent manifest isn't already loaded, so we'll
                          # just encode a fulltext of the manifest and pass that
                          # through to the revlog layer, and let it handle the delta
                          # process.
                          text = m.text()
                          arraytext = array.array('c', text)
                          cachedelta = None
                      n = self.addrevision(text, transaction, link, p1, p2, cachedelta)
                      self._mancache[n] = (m, arraytext)
                      return n

mercurial/match.py

0 +3 0

              # match.py - filename matching
              #
              #  Copyright 2008, 2009 Matt Mackall <mpm@selenic.com> and others
              #
              # This software may be used and distributed according to the terms of the
              # GNU General Public License version 2 or any later version.
              import re
              import util, pathutil
              from i18n import _
              def _rematcher(regex):
                  '''compile the regexp with the best available regexp engine and return a
                  matcher function'''
                  m = util.re.compile(regex)
                  try:
                      # slightly faster, provided by facebook's re2 bindings
                      return m.test_match
                  except AttributeError:
                      return m.match
              def _expandsets(kindpats, ctx):
                  '''Returns the kindpats list with the 'set' patterns expanded.'''
                  fset = set()
                  other = []
                  for kind, pat in kindpats:
                      if kind == 'set':
                          if not ctx:
                              raise util.Abort("fileset expression with no context")
                          s = ctx.getfileset(pat)
                          fset.update(s)
                          continue
                      other.append((kind, pat))
                  return fset, other
              def _kindpatsalwaysmatch(kindpats):
                  """"Checks whether the kindspats match everything, as e.g.
                  'relpath:.' does.
                  """
                  for kind, pat in kindpats:
                      if pat != '' or kind not in ['relpath', 'glob']:
                          return False
                  return True
              class match(object):
                  def __init__(self, root, cwd, patterns, include=[], exclude=[],
                               default='glob', exact=False, auditor=None, ctx=None):
                      """build an object to match a set of file patterns
                      arguments:
                      root - the canonical root of the tree you're matching against
                      cwd - the current working directory, if relevant
                      patterns - patterns to find
                      include - patterns to include (unless they are excluded)
                      exclude - patterns to exclude (even if they are included)
                      default - if a pattern in patterns has no explicit type, assume this one
                      exact - patterns are actually filenames (include/exclude still apply)
                      a pattern is one of:
                      'glob:<glob>' - a glob relative to cwd
                      're:<regexp>' - a regular expression
                      'path:<path>' - a path relative to repository root
                      'relglob:<glob>' - an unrooted glob (*.c matches C files in all dirs)
                      'relpath:<path>' - a path relative to cwd
                      'relre:<regexp>' - a regexp that needn't match the start of a name
                      'set:<fileset>' - a fileset expression
                      '<something>' - a pattern of the specified default type
                      """
                      self._root = root
                      self._cwd = cwd
                      self._files = [] # exact files and roots of patterns
                      self._anypats = bool(include or exclude)
                      self._ctx = ctx
                      self._always = False
                      self._pathrestricted = bool(include or exclude or patterns)
                      matchfns = []
                      if include:
                          kindpats = _normalize(include, 'glob', root, cwd, auditor)
                          self.includepat, im = _buildmatch(ctx, kindpats, '(?:/|$)')
                          matchfns.append(im)
                      if exclude:
                          kindpats = _normalize(exclude, 'glob', root, cwd, auditor)
                          self.excludepat, em = _buildmatch(ctx, kindpats, '(?:/|$)')
                          matchfns.append(lambda f: not em(f))
                      if exact:
                          if isinstance(patterns, list):
                              self._files = patterns
                          else:
                              self._files = list(patterns)
                          matchfns.append(self.exact)
                      elif patterns:
                          kindpats = _normalize(patterns, default, root, cwd, auditor)
                          if not _kindpatsalwaysmatch(kindpats):
                              self._files = _roots(kindpats)
                              self._anypats = self._anypats or _anypats(kindpats)
                              self.patternspat, pm = _buildmatch(ctx, kindpats, '$')
                              matchfns.append(pm)
                      if not matchfns:
                          m = util.always
                          self._always = True
                      elif len(matchfns) == 1:
                          m = matchfns[0]
                      else:
                          def m(f):
                              for matchfn in matchfns:
                                  if not matchfn(f):
                                      return False
                              return True
                      self.matchfn = m
                      self._fmap = set(self._files)
                  def __call__(self, fn):
                      return self.matchfn(fn)
                  def __iter__(self):
                      for f in self._files:
                          yield f
                  # Callbacks related to how the matcher is used by dirstate.walk.
                  # Subscribers to these events must monkeypatch the matcher object.
                  def bad(self, f, msg):
                      '''Callback from dirstate.walk for each explicit file that can't be
                      found/accessed, with an error message.'''
                      pass
                  # If an explicitdir is set, it will be called when an explicitly listed
                  # directory is visited.
                  explicitdir = None
                  # If an traversedir is set, it will be called when a directory discovered
                  # by recursive traversal is visited.
                  traversedir = None
                  def abs(self, f):
                      '''Convert a repo path back to path that is relative to the root of the
                      matcher.'''
                      return f
                  def rel(self, f):
                      '''Convert repo path back to path that is relative to cwd of matcher.'''
                      return util.pathto(self._root, self._cwd, f)
                  def uipath(self, f):
                      '''Convert repo path to a display path.  If patterns or -I/-X were used
                      to create this matcher, the display path will be relative to cwd.
                      Otherwise it is relative to the root of the repo.'''
                      return (self._pathrestricted and self.rel(f)) or self.abs(f)
                  def files(self):
                      '''Explicitly listed files or patterns or roots:
                      if no patterns or .always(): empty list,
                      if exact: list exact files,
                      if not .anypats(): list all files and dirs,
                      else: optimal roots'''
                      return self._files
                  def exact(self, f):
                      '''Returns True if f is in .files().'''
                      return f in self._fmap
                  def anypats(self):
                      '''Matcher uses patterns or include/exclude.'''
                      return self._anypats
                  def always(self):
                      '''Matcher will match everything and .files() will be empty
                      - optimization might be possible and necessary.'''
                      return self._always
+                 def isexact(self):
+                     return self.matchfn == self.exact
              def exact(root, cwd, files):
                  return match(root, cwd, files, exact=True)
              def always(root, cwd):
                  return match(root, cwd, [])
              class narrowmatcher(match):
                  """Adapt a matcher to work on a subdirectory only.
                  The paths are remapped to remove/insert the path as needed:
                  >>> m1 = match('root', '', ['a.txt', 'sub/b.txt'])
                  >>> m2 = narrowmatcher('sub', m1)
                  >>> bool(m2('a.txt'))
                  False
                  >>> bool(m2('b.txt'))
                  True
                  >>> bool(m2.matchfn('a.txt'))
                  False
                  >>> bool(m2.matchfn('b.txt'))
                  True
                  >>> m2.files()
                  ['b.txt']
                  >>> m2.exact('b.txt')
                  True
                  >>> util.pconvert(m2.rel('b.txt'))
                  'sub/b.txt'
                  >>> def bad(f, msg):
                  ...     print "%s: %s" % (f, msg)
                  >>> m1.bad = bad
                  >>> m2.bad('x.txt', 'No such file')
                  sub/x.txt: No such file
                  >>> m2.abs('c.txt')
                  'sub/c.txt'
                  """
                  def __init__(self, path, matcher):
                      self._root = matcher._root
                      self._cwd = matcher._cwd
                      self._path = path
                      self._matcher = matcher
                      self._always = matcher._always
                      self._pathrestricted = matcher._pathrestricted
                      self._files = [f[len(path) + 1:] for f in matcher._files
                                     if f.startswith(path + "/")]
                      self._anypats = matcher._anypats
                      self.matchfn = lambda fn: matcher.matchfn(self._path + "/" + fn)
                      self._fmap = set(self._files)
                  def abs(self, f):
                      return self._matcher.abs(self._path + "/" + f)
                  def bad(self, f, msg):
                      self._matcher.bad(self._path + "/" + f, msg)
                  def rel(self, f):
                      return self._matcher.rel(self._path + "/" + f)
              def patkind(pattern, default=None):
                  '''If pattern is 'kind:pat' with a known kind, return kind.'''
                  return _patsplit(pattern, default)[0]
              def _patsplit(pattern, default):
                  """Split a string into the optional pattern kind prefix and the actual
                  pattern."""
                  if ':' in pattern:
                      kind, pat = pattern.split(':', 1)
                      if kind in ('re', 'glob', 'path', 'relglob', 'relpath', 'relre',
                                  'listfile', 'listfile0', 'set'):
                          return kind, pat
                  return default, pattern
              def _globre(pat):
                  r'''Convert an extended glob string to a regexp string.
                  >>> print _globre(r'?')
                  .
                  >>> print _globre(r'*')
                  [^/]*
                  >>> print _globre(r'**')
                  .*
                  >>> print _globre(r'**/a')
                  (?:.*/)?a
                  >>> print _globre(r'a/**/b')
                  a\/(?:.*/)?b
                  >>> print _globre(r'[a*?!^][^b][!c]')
                  [a*?!^][\^b][^c]
                  >>> print _globre(r'{a,b}')
                  (?:a|b)
                  >>> print _globre(r'.\*\?')
                  \.\*\?
                  '''
                  i, n = 0, len(pat)
                  res = ''
                  group = 0
                  escape = util.re.escape
                  def peek():
                      return i < n and pat[i]
                  while i < n:
                      c = pat[i]
                      i += 1
                      if c not in '*?[{},\\':
                          res += escape(c)
                      elif c == '*':
                          if peek() == '*':
                              i += 1
                              if peek() == '/':
                                  i += 1
                                  res += '(?:.*/)?'
                              else:
                                  res += '.*'
                          else:
                              res += '[^/]*'
                      elif c == '?':
                          res += '.'
                      elif c == '[':
                          j = i
                          if j < n and pat[j] in '!]':
                              j += 1
                          while j < n and pat[j] != ']':
                              j += 1
                          if j >= n:
                              res += '\\['
                          else:
                              stuff = pat[i:j].replace('\\','\\\\')
                              i = j + 1
                              if stuff[0] == '!':
                                  stuff = '^' + stuff[1:]
                              elif stuff[0] == '^':
                                  stuff = '\\' + stuff
                              res = '%s[%s]' % (res, stuff)
                      elif c == '{':
                          group += 1
                          res += '(?:'
                      elif c == '}' and group:
                          res += ')'
                          group -= 1
                      elif c == ',' and group:
                          res += '|'
                      elif c == '\\':
                          p = peek()
                          if p:
                              i += 1
                              res += escape(p)
                          else:
                              res += escape(c)
                      else:
                          res += escape(c)
                  return res
              def _regex(kind, pat, globsuffix):
                  '''Convert a (normalized) pattern of any kind into a regular expression.
                  globsuffix is appended to the regexp of globs.'''
                  if not pat:
                      return ''
                  if kind == 're':
                      return pat
                  if kind == 'path':
                      return '^' + util.re.escape(pat) + '(?:/|$)'
                  if kind == 'relglob':
                      return '(?:|.*/)' + _globre(pat) + globsuffix
                  if kind == 'relpath':
                      return util.re.escape(pat) + '(?:/|$)'
                  if kind == 'relre':
                      if pat.startswith('^'):
                          return pat
                      return '.*' + pat
                  return _globre(pat) + globsuffix
              def _buildmatch(ctx, kindpats, globsuffix):
                  '''Return regexp string and a matcher function for kindpats.
                  globsuffix is appended to the regexp of globs.'''
                  fset, kindpats = _expandsets(kindpats, ctx)
                  if not kindpats:
                      return "", fset.__contains__
                  regex, mf = _buildregexmatch(kindpats, globsuffix)
                  if fset:
                      return regex, lambda f: f in fset or mf(f)
                  return regex, mf
              def _buildregexmatch(kindpats, globsuffix):
                  """Build a match function from a list of kinds and kindpats,
                  return regexp string and a matcher function."""
                  try:
                      regex = '(?:%s)' % '|'.join([_regex(k, p, globsuffix)
                                                   for (k, p) in kindpats])
                      if len(regex) > 20000:
                          raise OverflowError
                      return regex, _rematcher(regex)
                  except OverflowError:
                      # We're using a Python with a tiny regex engine and we
                      # made it explode, so we'll divide the pattern list in two
                      # until it works
                      l = len(kindpats)
                      if l < 2:
                          raise
                      regexa, a = _buildregexmatch(kindpats[:l//2], globsuffix)
                      regexb, b = _buildregexmatch(kindpats[l//2:], globsuffix)
                      return regex, lambda s: a(s) or b(s)
                  except re.error:
                      for k, p in kindpats:
                          try:
                              _rematcher('(?:%s)' % _regex(k, p, globsuffix))
                          except re.error:
                              raise util.Abort(_("invalid pattern (%s): %s") % (k, p))
                      raise util.Abort(_("invalid pattern"))
              def _normalize(patterns, default, root, cwd, auditor):
                  '''Convert 'kind:pat' from the patterns list to tuples with kind and
                  normalized and rooted patterns and with listfiles expanded.'''
                  kindpats = []
                  for kind, pat in [_patsplit(p, default) for p in patterns]:
                      if kind in ('glob', 'relpath'):
                          pat = pathutil.canonpath(root, cwd, pat, auditor)
                      elif kind in ('relglob', 'path'):
                          pat = util.normpath(pat)
                      elif kind in ('listfile', 'listfile0'):
                          try:
                              files = util.readfile(pat)
                              if kind == 'listfile0':
                                  files = files.split('\0')
                              else:
                                  files = files.splitlines()
                              files = [f for f in files if f]
                          except EnvironmentError:
                              raise util.Abort(_("unable to read file list (%s)") % pat)
                          kindpats += _normalize(files, default, root, cwd, auditor)
                          continue
                      # else: re or relre - which cannot be normalized
                      kindpats.append((kind, pat))
                  return kindpats
              def _roots(kindpats):
                  '''return roots and exact explicitly listed files from patterns
                  >>> _roots([('glob', 'g/*'), ('glob', 'g'), ('glob', 'g*')])
                  ['g', 'g', '.']
                  >>> _roots([('relpath', 'r'), ('path', 'p/p'), ('path', '')])
                  ['r', 'p/p', '.']
                  >>> _roots([('relglob', 'rg*'), ('re', 're/'), ('relre', 'rr')])
                  ['.', '.', '.']
                  '''
                  r = []
                  for kind, pat in kindpats:
                      if kind == 'glob': # find the non-glob prefix
                          root = []
                          for p in pat.split('/'):
                              if '[' in p or '{' in p or '*' in p or '?' in p:
                                  break
                              root.append(p)
                          r.append('/'.join(root) or '.')
                      elif kind in ('relpath', 'path'):
                          r.append(pat or '.')
                      else: # relglob, re, relre
                          r.append('.')
                  return r
              def _anypats(kindpats):
                  for kind, pat in kindpats:
                      if kind in ('glob', 're', 'relglob', 'relre', 'set'):
                          return True

General Comments 0

Write
Preview

You need to be logged in to leave comments. Login now

No TODOs yet

	Site-wide shortcuts
/	Use quick search box
g h	Goto home page
g g	Goto my private gists page
g G	Goto my public gists page
g 0-9	Goto bookmarked items from 0-9
n r	New repository page
n g	New gist page

	Repositories
g s	Goto summary page
g c	Goto changelog page
g f	Goto files page
g F	Goto files page with file search activated
g p	Goto pull requests page
g o	Goto repository settings
g O	Goto repository access permissions settings
t s	Toggle sidebar on some pages