upstream/mercurial-mirror Files · mercurial/lsprof.py

manifest: avoid corruption by dropping removed files with pure (issue5801)...

manifest: avoid corruption by dropping removed files with pure (issue5801) Previously, removed files would simply be marked by overwriting the first byte with NUL and dropping their entry in `self.position`. But no effort was made to ignore them when compacting the dictionary into text form. This allowed them to slip into the manifest revision, since the code seems to be trying to minimize the string operations by copying as large a chunk as possible. As part of this, compact() walks the existing text based on entries in the `positions` list, and consumed everything up to the next position entry. This typically resulted in a ValueError complaining about unsorted manifest entries. Sometimes it seems that files do get dropped in large repos- it seems to correspond to there being a new entry that would take the same slot. A much more trivial problem is that if the only changes were removals, `_compact()` didn't even run because `__delitem__` doesn't add anything to `self.extradata`. Now there's an explicit variable to flag this, both to allow `_compact()` to run, and to avoid searching the manifest in cases where there are no removals. In practice, this behavior was mostly obscured by the check in fastdelta() which takes a different path that explicitly drops removed files if there are fewer than 1000 changes. However, timeless has a repo where after rebasing tens of commits, a totally different path[1] is taken that bypasses the change count check and hits this problem. [1] https://www.mercurial-scm.org/repo/hg/file/2338bdea4474/mercurial/manifest.py#l1511

Gregory Szorc - - Load All Authors

File last commit:

r40238:56ea22fa default


                r42569:0546ead3

stable

Download file

             lsprof.py
        
                    127 lines
            
             | 4.1 KiB
            
                | text/x-python
            
             |
                PythonLexer
            
             / mercurial / lsprof.py
          
                    History
                
                 |
                  Annotation
                 | Raw
                 |Copy content
                 |Copy permalink

      from __future__ import absolute_import, print_function

      import _lsprof

      import sys

      Profiler = _lsprof.Profiler

      # PyPy doesn't expose profiler_entry from the module.

      profiler_entry = getattr(_lsprof, 'profiler_entry', None)

      __all__ = ['profile', 'Stats']

      def profile(f, *args, **kwds):

          """XXX docstring"""

          p = Profiler()

          p.enable(subcalls=True, builtins=True)

          try:

              f(*args, **kwds)

          finally:

              p.disable()

          return Stats(p.getstats())

      class Stats(object):

          """XXX docstring"""

          def __init__(self, data):

              self.data = data

          def sort(self, crit=r"inlinetime"):

              """XXX docstring"""

              # profiler_entries isn't defined when running under PyPy.

              if profiler_entry:

                  if crit not in profiler_entry.__dict__:

                      raise ValueError("Can't sort by %s" % crit)

              elif self.data and not getattr(self.data[0], crit, None):

                  raise ValueError("Can't sort by %s" % crit)

              self.data.sort(key=lambda x: getattr(x, crit), reverse=True)

              for e in self.data:

                  if e.calls:

                      e.calls.sort(key=lambda x: getattr(x, crit), reverse=True)

          def pprint(self, top=None, file=None, limit=None, climit=None):

              """XXX docstring"""

              if file is None:

                  file = sys.stdout

              d = self.data

              if top is not None:

                  d = d[:top]

              cols = "% 12d %12d %11.4f %11.4f   %s\n"

              hcols = "% 12s %12s %12s %12s %s\n"

              file.write(hcols % ("CallCount", "Recursive", "Total(s)",

                                  "Inline(s)", "module:lineno(function)"))

              count = 0

              for e in d:

                  file.write(cols % (e.callcount, e.reccallcount, e.totaltime,

                                     e.inlinetime, label(e.code)))

                  count += 1

                  if limit is not None and count == limit:

                      return

                  ccount = 0

                  if climit and e.calls:

                      for se in e.calls:

                          file.write(cols % (se.callcount, se.reccallcount,

                                             se.totaltime, se.inlinetime,

                                             "    %s" % label(se.code)))

                          count += 1

                          ccount += 1

                          if limit is not None and count == limit:

                              return

                          if climit is not None and ccount == climit:

                              break

          def freeze(self):

              """Replace all references to code objects with string

              descriptions; this makes it possible to pickle the instance."""

              # this code is probably rather ickier than it needs to be!

              for i in range(len(self.data)):

                  e = self.data[i]

                  if not isinstance(e.code, str):

                      self.data[i] = type(e)((label(e.code),) + e[1:])

                  if e.calls:

                      for j in range(len(e.calls)):

                          se = e.calls[j]

                          if not isinstance(se.code, str):

                              e.calls[j] = type(se)((label(se.code),) + se[1:])

      _fn2mod = {}

      def label(code):

          if isinstance(code, str):

              if sys.version_info.major >= 3:

                  code = code.encode('latin-1')

              return code

          try:

              mname = _fn2mod[code.co_filename]

          except KeyError:

              for k, v in list(sys.modules.iteritems()):

                  if v is None:

                      continue

                  if not isinstance(getattr(v, '__file__', None), str):

                      continue

                  if v.__file__.startswith(code.co_filename):

                      mname = _fn2mod[code.co_filename] = k

                      break

              else:

                  mname = _fn2mod[code.co_filename] = r'<%s>' % code.co_filename

          res = r'%s:%d(%s)' % (mname, code.co_firstlineno, code.co_name)

          if sys.version_info.major >= 3:

              res = res.encode('latin-1')

          return res

      if __name__ == '__main__':

          import os

          sys.argv = sys.argv[1:]

          if not sys.argv:

              print("usage: lsprof.py <script> <arguments...>", file=sys.stderr)

              sys.exit(2)

          sys.path.insert(0, os.path.abspath(os.path.dirname(sys.argv[0])))

          stats = profile(execfile, sys.argv[0], globals(), locals())

          stats.sort()

          stats.pprint()

	Site-wide shortcuts
/	Use quick search box
g h	Goto home page
g g	Goto my private gists page
g G	Goto my public gists page
g 0-9	Goto bookmarked items from 0-9
n r	New repository page
n g	New gist page

	Repositories
g s	Goto summary page
g c	Goto changelog page
g f	Goto files page
g F	Goto files page with file search activated
g p	Goto pull requests page
g o	Goto repository settings
g O	Goto repository access permissions settings
t s	Toggle sidebar on some pages

				from __future__ import absolute_import, print_function

				import _lsprof
				import sys

				Profiler = _lsprof.Profiler

				# PyPy doesn't expose profiler_entry from the module.
				profiler_entry = getattr(_lsprof, 'profiler_entry', None)

				__all__ = ['profile', 'Stats']

				def profile(f, args, *kwds):
				"""XXX docstring"""
				p = Profiler()
				p.enable(subcalls=True, builtins=True)
				try:
				f(args, *kwds)
				finally:
				p.disable()
				return Stats(p.getstats())


				class Stats(object):
				"""XXX docstring"""

				def __init__(self, data):
				self.data = data

				def sort(self, crit=r"inlinetime"):
				"""XXX docstring"""
				# profiler_entries isn't defined when running under PyPy.
				if profiler_entry:
				if crit not in profiler_entry.__dict__:
				raise ValueError("Can't sort by %s" % crit)
				elif self.data and not getattr(self.data[0], crit, None):
				raise ValueError("Can't sort by %s" % crit)

				self.data.sort(key=lambda x: getattr(x, crit), reverse=True)
				for e in self.data:
				if e.calls:
				e.calls.sort(key=lambda x: getattr(x, crit), reverse=True)

				def pprint(self, top=None, file=None, limit=None, climit=None):
				"""XXX docstring"""
				if file is None:
				file = sys.stdout
				d = self.data
				if top is not None:
				d = d[:top]
				cols = "% 12d %12d %11.4f %11.4f %s\n"
				hcols = "% 12s %12s %12s %12s %s\n"
				file.write(hcols % ("CallCount", "Recursive", "Total(s)",
				"Inline(s)", "module:lineno(function)"))
				count = 0
				for e in d:
				file.write(cols % (e.callcount, e.reccallcount, e.totaltime,
				e.inlinetime, label(e.code)))
				count += 1
				if limit is not None and count == limit:
				return
				ccount = 0
				if climit and e.calls:
				for se in e.calls:
				file.write(cols % (se.callcount, se.reccallcount,
				se.totaltime, se.inlinetime,
				" %s" % label(se.code)))
				count += 1
				ccount += 1
				if limit is not None and count == limit:
				return
				if climit is not None and ccount == climit:
				break

				def freeze(self):
				"""Replace all references to code objects with string
				descriptions; this makes it possible to pickle the instance."""

				# this code is probably rather ickier than it needs to be!
				for i in range(len(self.data)):
				e = self.data[i]
				if not isinstance(e.code, str):
				self.data[i] = type(e)((label(e.code),) + e[1:])
				if e.calls:
				for j in range(len(e.calls)):
				se = e.calls[j]
				if not isinstance(se.code, str):
				e.calls[j] = type(se)((label(se.code),) + se[1:])

				_fn2mod = {}

				def label(code):
				if isinstance(code, str):
				if sys.version_info.major >= 3:
				code = code.encode('latin-1')
				return code
				try:
				mname = _fn2mod[code.co_filename]
				except KeyError:
				for k, v in list(sys.modules.iteritems()):
				if v is None:
				continue
				if not isinstance(getattr(v, '__file__', None), str):
				continue
				if v.__file__.startswith(code.co_filename):
				mname = _fn2mod[code.co_filename] = k
				break
				else:
				mname = _fn2mod[code.co_filename] = r'<%s>' % code.co_filename

				res = r'%s:%d(%s)' % (mname, code.co_firstlineno, code.co_name)

				if sys.version_info.major >= 3:
				res = res.encode('latin-1')

				return res

				if __name__ == '__main__':
				import os
				sys.argv = sys.argv[1:]
				if not sys.argv:
				print("usage: lsprof.py <script> <arguments...>", file=sys.stderr)
				sys.exit(2)
				sys.path.insert(0, os.path.abspath(os.path.dirname(sys.argv[0])))
				stats = profile(execfile, sys.argv[0], globals(), locals())
				stats.sort()
				stats.pprint()