upstream/mercurial-mirror Files · mercurial/dirstateutils/docket.py

match: add `filepath:` pattern to match an exact filepath relative to the root...

match: add `filepath:` pattern to match an exact filepath relative to the root It's useful in certain automated workflows to make sure we recurse in directories whose name conflicts with files in other revisions. In addition it makes it possible to avoid building a potentially costly regex, improving performance when the set of files to match explicitly is large. The benchmark below are run in the following configuration : # data-env-vars.name = mozilla-central-2018-08-01-zstd-sparse-revlog # benchmark.name = files # benchmark.variants.rev = tip # benchmark.variants.files = all-list-filepath-sorted # bin-env-vars.hg.flavor = no-rust It also includes timings using the re2 engine (through the `google-re2` module) to show how much can be saved by just using a better regexp engine. Pattern time (seconds) time using re2 ----------------------------------------------------------- just "." 0.4 0.4 list of "filepath:…" 1.3 1.3 list of "path:…" 25.7 3.9 list of patterns 29.7 10.4 As you can see, Without re2, using "filepath:" instead of "path:" is a huge win. With re2, it is still about three times faster to not have to build the regex.

Gregory Szorc - - Load All Authors

File last commit:

r49801:642e31cb default


                r51588:1c31b343

default

Download file

             docket.py
        
                    70 lines
            
             | 2.2 KiB
            
                | text/x-python
            
             |
                PythonLexer
            
             / mercurial / dirstateutils / docket.py
          
                    History
                
                 |
                  Annotation
                 | Raw
                 |Copy content
                 |Copy permalink

      # dirstatedocket.py - docket file for dirstate-v2

      #

      # Copyright Mercurial Contributors

      #

      # This software may be used and distributed according to the terms of the

      # GNU General Public License version 2 or any later version.

      import struct

      from ..revlogutils import docket as docket_mod

      from . import v2

      V2_FORMAT_MARKER = b"dirstate-v2\n"

      # * 12 bytes: format marker

      # * 32 bytes: node ID of the working directory's first parent

      # * 32 bytes: node ID of the working directory's second parent

      # * {TREE_METADATA_SIZE} bytes: tree metadata, parsed separately

      # * 4 bytes: big-endian used size of the data file

      # * 1 byte: length of the data file's UUID

      # * variable: data file's UUID

      #

      # Node IDs are null-padded if shorter than 32 bytes.

      # A data file shorter than the specified used size is corrupted (truncated)

      HEADER = struct.Struct(

          ">{}s32s32s{}sLB".format(len(V2_FORMAT_MARKER), v2.TREE_METADATA_SIZE)

      )

      class DirstateDocket:

          data_filename_pattern = b'dirstate.%s'

          def __init__(self, parents, data_size, tree_metadata, uuid):

              self.parents = parents

              self.data_size = data_size

              self.tree_metadata = tree_metadata

              self.uuid = uuid

          @classmethod

          def with_new_uuid(cls, parents, data_size, tree_metadata):

              return cls(parents, data_size, tree_metadata, docket_mod.make_uid())

          @classmethod

          def parse(cls, data, nodeconstants):

              if not data:

                  parents = (nodeconstants.nullid, nodeconstants.nullid)

                  return cls(parents, 0, b'', None)

              marker, p1, p2, meta, data_size, uuid_size = HEADER.unpack_from(data)

              if marker != V2_FORMAT_MARKER:

                  raise ValueError("expected dirstate-v2 marker")

              uuid = data[HEADER.size : HEADER.size + uuid_size]

              p1 = p1[: nodeconstants.nodelen]

              p2 = p2[: nodeconstants.nodelen]

              return cls((p1, p2), data_size, meta, uuid)

          def serialize(self):

              p1, p2 = self.parents

              header = HEADER.pack(

                  V2_FORMAT_MARKER,

                  p1,

                  p2,

                  self.tree_metadata,

                  self.data_size,

                  len(self.uuid),

              )

              return header + self.uuid

          def data_filename(self):

              return self.data_filename_pattern % self.uuid

	Site-wide shortcuts
/	Use quick search box
g h	Goto home page
g g	Goto my private gists page
g G	Goto my public gists page
g 0-9	Goto bookmarked items from 0-9
n r	New repository page
n g	New gist page

	Repositories
g s	Goto summary page
g c	Goto changelog page
g f	Goto files page
g F	Goto files page with file search activated
g p	Goto pull requests page
g o	Goto repository settings
g O	Goto repository access permissions settings
t s	Toggle sidebar on some pages

				# dirstatedocket.py - docket file for dirstate-v2
				#
				# Copyright Mercurial Contributors
				#
				# This software may be used and distributed according to the terms of the
				# GNU General Public License version 2 or any later version.


				import struct

				from ..revlogutils import docket as docket_mod
				from . import v2

				V2_FORMAT_MARKER = b"dirstate-v2\n"

				# * 12 bytes: format marker
				# * 32 bytes: node ID of the working directory's first parent
				# * 32 bytes: node ID of the working directory's second parent
				# * {TREE_METADATA_SIZE} bytes: tree metadata, parsed separately
				# * 4 bytes: big-endian used size of the data file
				# * 1 byte: length of the data file's UUID
				# * variable: data file's UUID
				#
				# Node IDs are null-padded if shorter than 32 bytes.
				# A data file shorter than the specified used size is corrupted (truncated)
				HEADER = struct.Struct(
				">{}s32s32s{}sLB".format(len(V2_FORMAT_MARKER), v2.TREE_METADATA_SIZE)
				)


				class DirstateDocket:
				data_filename_pattern = b'dirstate.%s'

				def __init__(self, parents, data_size, tree_metadata, uuid):
				self.parents = parents
				self.data_size = data_size
				self.tree_metadata = tree_metadata
				self.uuid = uuid

				@classmethod
				def with_new_uuid(cls, parents, data_size, tree_metadata):
				return cls(parents, data_size, tree_metadata, docket_mod.make_uid())

				@classmethod
				def parse(cls, data, nodeconstants):
				if not data:
				parents = (nodeconstants.nullid, nodeconstants.nullid)
				return cls(parents, 0, b'', None)
				marker, p1, p2, meta, data_size, uuid_size = HEADER.unpack_from(data)
				if marker != V2_FORMAT_MARKER:
				raise ValueError("expected dirstate-v2 marker")
				uuid = data[HEADER.size : HEADER.size + uuid_size]
				p1 = p1[: nodeconstants.nodelen]
				p2 = p2[: nodeconstants.nodelen]
				return cls((p1, p2), data_size, meta, uuid)

				def serialize(self):
				p1, p2 = self.parents
				header = HEADER.pack(
				V2_FORMAT_MARKER,
				p1,
				p2,
				self.tree_metadata,
				self.data_size,
				len(self.uuid),
				)
				return header + self.uuid

				def data_filename(self):
				return self.data_filename_pattern % self.uuid