upstream/mercurial-mirror Files · mercurial/filelog.py

localrepo.commit: normalize commit message even for rawcommit....

localrepo.commit: normalize commit message even for rawcommit. This normalization consists of: - stripping trailing whitespace - always using "\n" as the line separator I think the main reason rawcommit was skipping this normalization was an attempt to preserve hashes during an hg->hg conversion. While this is a nice goal, it's not particularly interesting in practice. Since SHA-1 is so strong, the only safe way to do it is to have absolutely identical revisions. But: - if the original revision was created with a recent version of hg, the commit message will be the same, with or without that normalization - if it was created with an ancient version of hg that didn't do any normalization, even if the commit message is identical, the file list in the changelog is likely to be different (e.g. no removed files), and there were some old issues with e.g. extra file merging, which will end up changing the hash anyway - in any case, if one *really* has to preserve hashes, it's easier (and faster) to fake a partial conversion using something like: hg clone -U -r rev orig-repo new-repo hg -R new-repo log --template '#node# #node#\n' > new-repo/.hg/shamap Additionally, we've had some reports of problems arising from this lack of normalization - e.g. issue871, and a user that was wondering why hg export/hg import was not preserving hashes when there was nothing unusual going on (it was just import doing the normalization that had been skipped). This also means that it's even more unlikely to get identical revisions when going $VCS->hg->$VCS.

Joel Rosdahl - - Load All Authors

File last commit:

r6212:e75aab65 default


                r6254:3667b6e4

default

Download file

             filelog.py
        
                    83 lines
            
             | 2.5 KiB
            
                | text/x-python
            
             |
                PythonLexer
            
             / mercurial / filelog.py
          
                    History
                
                 |
                  Annotation
                 | Raw
                 |Copy content
                 |Copy permalink

      # filelog.py - file history class for mercurial

      #

      # Copyright 2005-2007 Matt Mackall <mpm@selenic.com>

      #

      # This software may be used and distributed according to the terms

      # of the GNU General Public License, incorporated herein by reference.

      from node import bin, nullid

      from revlog import revlog

      class filelog(revlog):

          def __init__(self, opener, path):

              revlog.__init__(self, opener,

                              "/".join(("data", self.encodedir(path + ".i"))))

          # This avoids a collision between a file named foo and a dir named

          # foo.i or foo.d

          def encodedir(self, path):

              return (path

                      .replace(".hg/", ".hg.hg/")

                      .replace(".i/", ".i.hg/")

                      .replace(".d/", ".d.hg/"))

          def decodedir(self, path):

              return (path

                      .replace(".d.hg/", ".d/")

                      .replace(".i.hg/", ".i/")

                      .replace(".hg.hg/", ".hg/"))

          def read(self, node):

              t = self.revision(node)

              if not t.startswith('\1\n'):

                  return t

              s = t.index('\1\n', 2)

              return t[s+2:]

          def _readmeta(self, node):

              t = self.revision(node)

              if not t.startswith('\1\n'):

                  return {}

              s = t.index('\1\n', 2)

              mt = t[2:s]

              m = {}

              for l in mt.splitlines():

                  k, v = l.split(": ", 1)

                  m[k] = v

              return m

          def add(self, text, meta, transaction, link, p1=None, p2=None):

              if meta or text.startswith('\1\n'):

                  mt = ""

                  if meta:

                      mt = [ "%s: %s\n" % (k, v) for k,v in meta.items() ]

                  text = "\1\n%s\1\n%s" % ("".join(mt), text)

              return self.addrevision(text, transaction, link, p1, p2)

          def renamed(self, node):

              if self.parents(node)[0] != nullid:

                  return False

              m = self._readmeta(node)

              if m and "copy" in m:

                  return (m["copy"], bin(m["copyrev"]))

              return False

          def size(self, rev):

              """return the size of a given revision"""

              # for revisions with renames, we have to go the slow way

              node = self.node(rev)

              if self.renamed(node):

                  return len(self.read(node))

              return revlog.size(self, rev)

          def cmp(self, node, text):

              """compare text with a given file revision"""

              # for renames, we have to go the slow way

              if self.renamed(node):

                  t2 = self.read(node)

                  return t2 != text

              return revlog.cmp(self, node, text)

	Site-wide shortcuts
/	Use quick search box
g h	Goto home page
g g	Goto my private gists page
g G	Goto my public gists page
g 0-9	Goto bookmarked items from 0-9
n r	New repository page
n g	New gist page

	Repositories
g s	Goto summary page
g c	Goto changelog page
g f	Goto files page
g F	Goto files page with file search activated
g p	Goto pull requests page
g o	Goto repository settings
g O	Goto repository access permissions settings
t s	Toggle sidebar on some pages

				# filelog.py - file history class for mercurial
				#
				# Copyright 2005-2007 Matt Mackall <mpm@selenic.com>
				#
				# This software may be used and distributed according to the terms
				# of the GNU General Public License, incorporated herein by reference.

				from node import bin, nullid
				from revlog import revlog

				class filelog(revlog):
				def __init__(self, opener, path):
				revlog.__init__(self, opener,
				"/".join(("data", self.encodedir(path + ".i"))))

				# This avoids a collision between a file named foo and a dir named
				# foo.i or foo.d
				def encodedir(self, path):
				return (path
				.replace(".hg/", ".hg.hg/")
				.replace(".i/", ".i.hg/")
				.replace(".d/", ".d.hg/"))

				def decodedir(self, path):
				return (path
				.replace(".d.hg/", ".d/")
				.replace(".i.hg/", ".i/")
				.replace(".hg.hg/", ".hg/"))

				def read(self, node):
				t = self.revision(node)
				if not t.startswith('\1\n'):
				return t
				s = t.index('\1\n', 2)
				return t[s+2:]

				def _readmeta(self, node):
				t = self.revision(node)
				if not t.startswith('\1\n'):
				return {}
				s = t.index('\1\n', 2)
				mt = t[2:s]
				m = {}
				for l in mt.splitlines():
				k, v = l.split(": ", 1)
				m[k] = v
				return m

				def add(self, text, meta, transaction, link, p1=None, p2=None):
				if meta or text.startswith('\1\n'):
				mt = ""
				if meta:
				mt = [ "%s: %s\n" % (k, v) for k,v in meta.items() ]
				text = "\1\n%s\1\n%s" % ("".join(mt), text)
				return self.addrevision(text, transaction, link, p1, p2)

				def renamed(self, node):
				if self.parents(node)[0] != nullid:
				return False
				m = self._readmeta(node)
				if m and "copy" in m:
				return (m["copy"], bin(m["copyrev"]))
				return False

				def size(self, rev):
				"""return the size of a given revision"""

				# for revisions with renames, we have to go the slow way
				node = self.node(rev)
				if self.renamed(node):
				return len(self.read(node))

				return revlog.size(self, rev)

				def cmp(self, node, text):
				"""compare text with a given file revision"""

				# for renames, we have to go the slow way
				if self.renamed(node):
				t2 = self.read(node)
				return t2 != text

				return revlog.cmp(self, node, text)