upstream/mercurial-mirror Files · mercurial/filelog.py

patch: fix patch hunk/metdata synchronization (issue3384)...

patch: fix patch hunk/metdata synchronization (issue3384) Git patches are parsed in two phases: 1) extract metadata, 2) parse actual deltas and merge them with the previous metadata. We do this to avoid dependency issues like "modify a; copy a to b", where "b" must be copied from the unmodified "a". Issue3384 is caused by flaky code I wrote to synchronize the patch metadata with the emitted hunk: if (gitpatches and (gitpatches[-1][0] == afile or gitpatches[-1][1] == bfile)): gp = gitpatches.pop()[2] With a patch like: diff --git a/a b/c copy from a copy to c --- a/a +++ b/c @@ -1,1 +1,2 @@ a +a @@ -2,1 +2,2 @@ a +a diff --git a/a b/a --- a/a +++ b/a @@ -1,1 +1,2 @@ a +b the first hunk of the first block is matched with the metadata for the block "diff --git a/a b/c", then the second hunk of the first block is matched with the metadata of the second block "diff --git a/a b/a", because of the "or" in the code paste above. Turning the "or" into an "and" is not enough as we have to deal with /dev/null cases for each file. We I remove this broken piece of code: # copy/rename + modify should modify target, not source if gp.op in ('COPY', 'DELETE', 'RENAME', 'ADD') or gp.mode: afile = bfile because "afile = bfile" set "afile" to stuff like "b/file" instead of "a/file", and because this only happens for git patches, which afile/bfile are ignored anyway by applydiff(). v2: - Avoid a traceback on git metadata desynchronization

Sune Foldager - - Load All Authors

File last commit:

r14287:7c231754 default


                r16506:fc4e0fec

stable

Download file

             filelog.py
        
                    92 lines
            
             | 2.7 KiB
            
                | text/x-python
            
             |
                PythonLexer
            
             / mercurial / filelog.py
          
                    History
                
                 |
                  Annotation
                 | Raw
                 |Copy content
                 |Copy permalink

      # filelog.py - file history class for mercurial

      #

      # Copyright 2005-2007 Matt Mackall <mpm@selenic.com>

      #

      # This software may be used and distributed according to the terms of the

      # GNU General Public License version 2 or any later version.

      import revlog

      import re

      _mdre = re.compile('\1\n')

      def _parsemeta(text):

          """return (metadatadict, keylist, metadatasize)"""

          # text can be buffer, so we can't use .startswith or .index

          if text[:2] != '\1\n':

              return None, None, None

          s = _mdre.search(text, 2).start()

          mtext = text[2:s]

          meta = {}

          keys = []

          for l in mtext.splitlines():

              k, v = l.split(": ", 1)

              meta[k] = v

              keys.append(k)

          return meta, keys, (s + 2)

      def _packmeta(meta, keys=None):

          if not keys:

              keys = sorted(meta.iterkeys())

          return "".join("%s: %s\n" % (k, meta[k]) for k in keys)

      class filelog(revlog.revlog):

          def __init__(self, opener, path):

              revlog.revlog.__init__(self, opener,

                              "/".join(("data", path + ".i")))

          def read(self, node):

              t = self.revision(node)

              if not t.startswith('\1\n'):

                  return t

              s = t.index('\1\n', 2)

              return t[s + 2:]

          def add(self, text, meta, transaction, link, p1=None, p2=None):

              if meta or text.startswith('\1\n'):

                  text = "\1\n%s\1\n%s" % (_packmeta(meta), text)

              return self.addrevision(text, transaction, link, p1, p2)

          def renamed(self, node):

              if self.parents(node)[0] != revlog.nullid:

                  return False

              t = self.revision(node)

              m = _parsemeta(t)[0]

              if m and "copy" in m:

                  return (m["copy"], revlog.bin(m["copyrev"]))

              return False

          def size(self, rev):

              """return the size of a given revision"""

              # for revisions with renames, we have to go the slow way

              node = self.node(rev)

              if self.renamed(node):

                  return len(self.read(node))

              # XXX if self.read(node).startswith("\1\n"), this returns (size+4)

              return revlog.revlog.size(self, rev)

          def cmp(self, node, text):

              """compare text with a given file revision

              returns True if text is different than what is stored.

              """

              t = text

              if text.startswith('\1\n'):

                  t = '\1\n\1\n' + text

              samehashes = not revlog.revlog.cmp(self, node, t)

              if samehashes:

                  return False

              # renaming a file produces a different hash, even if the data

              # remains unchanged. Check if it's the case (slow):

              if self.renamed(node):

                  t2 = self.read(node)

                  return t2 != text

              return True

          def _file(self, f):

              return filelog(self.opener, f)

	Site-wide shortcuts
/	Use quick search box
g h	Goto home page
g g	Goto my private gists page
g G	Goto my public gists page
g 0-9	Goto bookmarked items from 0-9
n r	New repository page
n g	New gist page

	Repositories
g s	Goto summary page
g c	Goto changelog page
g f	Goto files page
g F	Goto files page with file search activated
g p	Goto pull requests page
g o	Goto repository settings
g O	Goto repository access permissions settings
t s	Toggle sidebar on some pages

				# filelog.py - file history class for mercurial
				#
				# Copyright 2005-2007 Matt Mackall <mpm@selenic.com>
				#
				# This software may be used and distributed according to the terms of the
				# GNU General Public License version 2 or any later version.

				import revlog
				import re

				_mdre = re.compile('\1\n')
				def _parsemeta(text):
				"""return (metadatadict, keylist, metadatasize)"""
				# text can be buffer, so we can't use .startswith or .index
				if text[:2] != '\1\n':
				return None, None, None
				s = _mdre.search(text, 2).start()
				mtext = text[2:s]
				meta = {}
				keys = []
				for l in mtext.splitlines():
				k, v = l.split(": ", 1)
				meta[k] = v
				keys.append(k)
				return meta, keys, (s + 2)

				def _packmeta(meta, keys=None):
				if not keys:
				keys = sorted(meta.iterkeys())
				return "".join("%s: %s\n" % (k, meta[k]) for k in keys)

				class filelog(revlog.revlog):
				def __init__(self, opener, path):
				revlog.revlog.__init__(self, opener,
				"/".join(("data", path + ".i")))

				def read(self, node):
				t = self.revision(node)
				if not t.startswith('\1\n'):
				return t
				s = t.index('\1\n', 2)
				return t[s + 2:]

				def add(self, text, meta, transaction, link, p1=None, p2=None):
				if meta or text.startswith('\1\n'):
				text = "\1\n%s\1\n%s" % (_packmeta(meta), text)
				return self.addrevision(text, transaction, link, p1, p2)

				def renamed(self, node):
				if self.parents(node)[0] != revlog.nullid:
				return False
				t = self.revision(node)
				m = _parsemeta(t)[0]
				if m and "copy" in m:
				return (m["copy"], revlog.bin(m["copyrev"]))
				return False

				def size(self, rev):
				"""return the size of a given revision"""

				# for revisions with renames, we have to go the slow way
				node = self.node(rev)
				if self.renamed(node):
				return len(self.read(node))

				# XXX if self.read(node).startswith("\1\n"), this returns (size+4)
				return revlog.revlog.size(self, rev)

				def cmp(self, node, text):
				"""compare text with a given file revision

				returns True if text is different than what is stored.
				"""

				t = text
				if text.startswith('\1\n'):
				t = '\1\n\1\n' + text

				samehashes = not revlog.revlog.cmp(self, node, t)
				if samehashes:
				return False

				# renaming a file produces a different hash, even if the data
				# remains unchanged. Check if it's the case (slow):
				if self.renamed(node):
				t2 = self.read(node)
				return t2 != text

				return True

				def _file(self, f):
				return filelog(self.opener, f)