upstream/mercurial-mirror Files · mercurial/repair.py

spanset: directly use __contains__ instead of a lambda...

spanset: directly use __contains__ instead of a lambda Spanset are massively used in revset. First because the initial subset itself is a repo wide spanset. We speed up the __and__ operation by getting rid of a gratuitous lambda call. A more long terms solution would be to: 1. speed up operation between spansets, 2. have a special smartset for `all` revisions. In the mean time, this is a very simple fix that buyback some of the performance regression. Below is performance benchmark for trival `and` operation between two spansets. (Run on an unspecified fairly large repository.) revset tip:0 2.9.2) wall 0.282543 comb 0.280000 user 0.260000 sys 0.020000 (best of 35) before) wall 0.819181 comb 0.820000 user 0.820000 sys 0.000000 (best of 12) after) wall 0.645358 comb 0.650000 user 0.650000 sys 0.000000 (best of 16) Proof of concept implementation of an `all` smartset brings this to 0.10 but it's too invasive for stable.

Pierre-Yves David - - Load All Authors

File last commit:

r21064:4d9d490d default


                r21207:b9defeeb

stable

Download file

             repair.py
        
                    183 lines
            
             | 6.1 KiB
            
                | text/x-python
            
             |
                PythonLexer
            
             / mercurial / repair.py
          
                    History
                
                 |
                  Annotation
                 | Raw
                 |Copy content
                 |Copy permalink

      # repair.py - functions for repository repair for mercurial

      #

      # Copyright 2005, 2006 Chris Mason <mason@suse.com>

      # Copyright 2007 Matt Mackall

      #

      # This software may be used and distributed according to the terms of the

      # GNU General Public License version 2 or any later version.

      from mercurial import changegroup, exchange

      from mercurial.node import short

      from mercurial.i18n import _

      import errno

      def _bundle(repo, bases, heads, node, suffix, compress=True):

          """create a bundle with the specified revisions as a backup"""

          cg = changegroup.changegroupsubset(repo, bases, heads, 'strip')

          backupdir = "strip-backup"

          vfs = repo.vfs

          if not vfs.isdir(backupdir):

              vfs.mkdir(backupdir)

          name = "%s/%s-%s.hg" % (backupdir, short(node), suffix)

          if compress:

              bundletype = "HG10BZ"

          else:

              bundletype = "HG10UN"

          return changegroup.writebundle(cg, name, bundletype, vfs)

      def _collectfiles(repo, striprev):

          """find out the filelogs affected by the strip"""

          files = set()

          for x in xrange(striprev, len(repo)):

              files.update(repo[x].files())

          return sorted(files)

      def _collectbrokencsets(repo, files, striprev):

          """return the changesets which will be broken by the truncation"""

          s = set()

          def collectone(revlog):

              _, brokenset = revlog.getstrippoint(striprev)

              s.update([revlog.linkrev(r) for r in brokenset])

          collectone(repo.manifest)

          for fname in files:

              collectone(repo.file(fname))

          return s

      def strip(ui, repo, nodelist, backup="all", topic='backup'):

          repo = repo.unfiltered()

          repo.destroying()

          cl = repo.changelog

          # TODO handle undo of merge sets

          if isinstance(nodelist, str):

              nodelist = [nodelist]

          striplist = [cl.rev(node) for node in nodelist]

          striprev = min(striplist)

          keeppartialbundle = backup == 'strip'

          # Some revisions with rev > striprev may not be descendants of striprev.

          # We have to find these revisions and put them in a bundle, so that

          # we can restore them after the truncations.

          # To create the bundle we use repo.changegroupsubset which requires

          # the list of heads and bases of the set of interesting revisions.

          # (head = revision in the set that has no descendant in the set;

          #  base = revision in the set that has no ancestor in the set)

          tostrip = set(striplist)

          for rev in striplist:

              for desc in cl.descendants([rev]):

                  tostrip.add(desc)

          files = _collectfiles(repo, striprev)

          saverevs = _collectbrokencsets(repo, files, striprev)

          # compute heads

          saveheads = set(saverevs)

          for r in xrange(striprev + 1, len(cl)):

              if r not in tostrip:

                  saverevs.add(r)

                  saveheads.difference_update(cl.parentrevs(r))

                  saveheads.add(r)

          saveheads = [cl.node(r) for r in saveheads]

          # compute base nodes

          if saverevs:

              descendants = set(cl.descendants(saverevs))

              saverevs.difference_update(descendants)

          savebases = [cl.node(r) for r in saverevs]

          stripbases = [cl.node(r) for r in tostrip]

          # For a set s, max(parents(s) - s) is the same as max(heads(::s - s)), but

          # is much faster

          newbmtarget = repo.revs('max(parents(%ld) - (%ld))', tostrip, tostrip)

          if newbmtarget:

              newbmtarget = repo[newbmtarget[0]].node()

          else:

              newbmtarget = '.'

          bm = repo._bookmarks

          updatebm = []

          for m in bm:

              rev = repo[bm[m]].rev()

              if rev in tostrip:

                  updatebm.append(m)

          # create a changegroup for all the branches we need to keep

          backupfile = None

          vfs = repo.vfs

          if backup == "all":

              backupfile = _bundle(repo, stripbases, cl.heads(), node, topic)

              repo.ui.status(_("saved backup bundle to %s\n") %

                             vfs.join(backupfile))

              repo.ui.log("backupbundle", "saved backup bundle to %s\n",

                          vfs.join(backupfile))

          if saveheads or savebases:

              # do not compress partial bundle if we remove it from disk later

              chgrpfile = _bundle(repo, savebases, saveheads, node, 'temp',

                                  compress=keeppartialbundle)

          mfst = repo.manifest

          tr = repo.transaction("strip")

          offset = len(tr.entries)

          try:

              tr.startgroup()

              cl.strip(striprev, tr)

              mfst.strip(striprev, tr)

              for fn in files:

                  repo.file(fn).strip(striprev, tr)

              tr.endgroup()

              try:

                  for i in xrange(offset, len(tr.entries)):

                      file, troffset, ignore = tr.entries[i]

                      repo.sopener(file, 'a').truncate(troffset)

                      if troffset == 0:

                          repo.store.markremoved(file)

                  tr.close()

              except: # re-raises

                  tr.abort()

                  raise

              if saveheads or savebases:

                  ui.note(_("adding branch\n"))

                  f = vfs.open(chgrpfile, "rb")

                  gen = exchange.readbundle(ui, f, chgrpfile, vfs)

                  if not repo.ui.verbose:

                      # silence internal shuffling chatter

                      repo.ui.pushbuffer()

                  changegroup.addchangegroup(repo, gen, 'strip',

                                             'bundle:' + vfs.join(chgrpfile), True)

                  if not repo.ui.verbose:

                      repo.ui.popbuffer()

                  f.close()

                  if not keeppartialbundle:

                      vfs.unlink(chgrpfile)

              # remove undo files

              for undovfs, undofile in repo.undofiles():

                  try:

                      undovfs.unlink(undofile)

                  except OSError, e:

                      if e.errno != errno.ENOENT:

                          ui.warn(_('error removing %s: %s\n') %

                                  (undovfs.join(undofile), str(e)))

              for m in updatebm:

                  bm[m] = repo[newbmtarget].node()

              bm.write()

          except: # re-raises

              if backupfile:

                  ui.warn(_("strip failed, full bundle stored in '%s'\n")

                          % vfs.join(backupfile))

              elif saveheads:

                  ui.warn(_("strip failed, partial bundle stored in '%s'\n")

                          % vfs.join(chgrpfile))

              raise

          repo.destroyed()

	Site-wide shortcuts
/	Use quick search box
g h	Goto home page
g g	Goto my private gists page
g G	Goto my public gists page
g 0-9	Goto bookmarked items from 0-9
n r	New repository page
n g	New gist page

	Repositories
g s	Goto summary page
g c	Goto changelog page
g f	Goto files page
g F	Goto files page with file search activated
g p	Goto pull requests page
g o	Goto repository settings
g O	Goto repository access permissions settings
t s	Toggle sidebar on some pages

				# repair.py - functions for repository repair for mercurial
				#
				# Copyright 2005, 2006 Chris Mason <mason@suse.com>
				# Copyright 2007 Matt Mackall
				#
				# This software may be used and distributed according to the terms of the
				# GNU General Public License version 2 or any later version.

				from mercurial import changegroup, exchange
				from mercurial.node import short
				from mercurial.i18n import _
				import errno

				def _bundle(repo, bases, heads, node, suffix, compress=True):
				"""create a bundle with the specified revisions as a backup"""
				cg = changegroup.changegroupsubset(repo, bases, heads, 'strip')
				backupdir = "strip-backup"
				vfs = repo.vfs
				if not vfs.isdir(backupdir):
				vfs.mkdir(backupdir)
				name = "%s/%s-%s.hg" % (backupdir, short(node), suffix)
				if compress:
				bundletype = "HG10BZ"
				else:
				bundletype = "HG10UN"
				return changegroup.writebundle(cg, name, bundletype, vfs)

				def _collectfiles(repo, striprev):
				"""find out the filelogs affected by the strip"""
				files = set()

				for x in xrange(striprev, len(repo)):
				files.update(repo[x].files())

				return sorted(files)

				def _collectbrokencsets(repo, files, striprev):
				"""return the changesets which will be broken by the truncation"""
				s = set()
				def collectone(revlog):
				_, brokenset = revlog.getstrippoint(striprev)
				s.update([revlog.linkrev(r) for r in brokenset])

				collectone(repo.manifest)
				for fname in files:
				collectone(repo.file(fname))

				return s

				def strip(ui, repo, nodelist, backup="all", topic='backup'):
				repo = repo.unfiltered()
				repo.destroying()

				cl = repo.changelog
				# TODO handle undo of merge sets
				if isinstance(nodelist, str):
				nodelist = [nodelist]
				striplist = [cl.rev(node) for node in nodelist]
				striprev = min(striplist)

				keeppartialbundle = backup == 'strip'

				# Some revisions with rev > striprev may not be descendants of striprev.
				# We have to find these revisions and put them in a bundle, so that
				# we can restore them after the truncations.
				# To create the bundle we use repo.changegroupsubset which requires
				# the list of heads and bases of the set of interesting revisions.
				# (head = revision in the set that has no descendant in the set;
				# base = revision in the set that has no ancestor in the set)
				tostrip = set(striplist)
				for rev in striplist:
				for desc in cl.descendants([rev]):
				tostrip.add(desc)

				files = _collectfiles(repo, striprev)
				saverevs = _collectbrokencsets(repo, files, striprev)

				# compute heads
				saveheads = set(saverevs)
				for r in xrange(striprev + 1, len(cl)):
				if r not in tostrip:
				saverevs.add(r)
				saveheads.difference_update(cl.parentrevs(r))
				saveheads.add(r)
				saveheads = [cl.node(r) for r in saveheads]

				# compute base nodes
				if saverevs:
				descendants = set(cl.descendants(saverevs))
				saverevs.difference_update(descendants)
				savebases = [cl.node(r) for r in saverevs]
				stripbases = [cl.node(r) for r in tostrip]

				# For a set s, max(parents(s) - s) is the same as max(heads(::s - s)), but
				# is much faster
				newbmtarget = repo.revs('max(parents(%ld) - (%ld))', tostrip, tostrip)
				if newbmtarget:
				newbmtarget = repo[newbmtarget[0]].node()
				else:
				newbmtarget = '.'

				bm = repo._bookmarks
				updatebm = []
				for m in bm:
				rev = repo[bm[m]].rev()
				if rev in tostrip:
				updatebm.append(m)

				# create a changegroup for all the branches we need to keep
				backupfile = None
				vfs = repo.vfs
				if backup == "all":
				backupfile = _bundle(repo, stripbases, cl.heads(), node, topic)
				repo.ui.status(_("saved backup bundle to %s\n") %
				vfs.join(backupfile))
				repo.ui.log("backupbundle", "saved backup bundle to %s\n",
				vfs.join(backupfile))
				if saveheads or savebases:
				# do not compress partial bundle if we remove it from disk later
				chgrpfile = _bundle(repo, savebases, saveheads, node, 'temp',
				compress=keeppartialbundle)

				mfst = repo.manifest

				tr = repo.transaction("strip")
				offset = len(tr.entries)

				try:
				tr.startgroup()
				cl.strip(striprev, tr)
				mfst.strip(striprev, tr)
				for fn in files:
				repo.file(fn).strip(striprev, tr)
				tr.endgroup()

				try:
				for i in xrange(offset, len(tr.entries)):
				file, troffset, ignore = tr.entries[i]
				repo.sopener(file, 'a').truncate(troffset)
				if troffset == 0:
				repo.store.markremoved(file)
				tr.close()
				except: # re-raises
				tr.abort()
				raise

				if saveheads or savebases:
				ui.note(_("adding branch\n"))
				f = vfs.open(chgrpfile, "rb")
				gen = exchange.readbundle(ui, f, chgrpfile, vfs)
				if not repo.ui.verbose:
				# silence internal shuffling chatter
				repo.ui.pushbuffer()
				changegroup.addchangegroup(repo, gen, 'strip',
				'bundle:' + vfs.join(chgrpfile), True)
				if not repo.ui.verbose:
				repo.ui.popbuffer()
				f.close()
				if not keeppartialbundle:
				vfs.unlink(chgrpfile)

				# remove undo files
				for undovfs, undofile in repo.undofiles():
				try:
				undovfs.unlink(undofile)
				except OSError, e:
				if e.errno != errno.ENOENT:
				ui.warn(_('error removing %s: %s\n') %
				(undovfs.join(undofile), str(e)))

				for m in updatebm:
				bm[m] = repo[newbmtarget].node()
				bm.write()
				except: # re-raises
				if backupfile:
				ui.warn(_("strip failed, full bundle stored in '%s'\n")
				% vfs.join(backupfile))
				elif saveheads:
				ui.warn(_("strip failed, partial bundle stored in '%s'\n")
				% vfs.join(chgrpfile))
				raise

				repo.destroyed()