changegroup.py
153 lines
| 4.7 KiB
| text/x-python
|
PythonLexer
/ mercurial / changegroup.py
Martin Geisler
|
r8226 | # changegroup.py - Mercurial changegroup manipulation functions | ||
# | ||||
# Copyright 2006 Matt Mackall <mpm@selenic.com> | ||||
# | ||||
# This software may be used and distributed according to the terms of the | ||||
Matt Mackall
|
r10263 | # GNU General Public License version 2 or any later version. | ||
Matt Mackall
|
r3877 | |||
Matt Mackall
|
r3891 | from i18n import _ | ||
Simon Heimberg
|
r8312 | import util | ||
import struct, os, bz2, zlib, tempfile | ||||
Thomas Arendsen Hein
|
r1981 | |||
def getchunk(source): | ||||
Greg Ward
|
r9437 | """return the next chunk from changegroup 'source' as a string""" | ||
Thomas Arendsen Hein
|
r1981 | d = source.read(4) | ||
if not d: | ||||
return "" | ||||
l = struct.unpack(">l", d)[0] | ||||
if l <= 4: | ||||
return "" | ||||
d = source.read(l - 4) | ||||
if len(d) < l - 4: | ||||
raise util.Abort(_("premature EOF reading chunk" | ||||
" (got %d bytes, expected %d)") | ||||
% (len(d), l - 4)) | ||||
return d | ||||
Augie Fackler
|
r10430 | def chunkiter(source, progress=None): | ||
Greg Ward
|
r9437 | """iterate through the chunks in source, yielding a sequence of chunks | ||
(strings)""" | ||||
Thomas Arendsen Hein
|
r1981 | while 1: | ||
c = getchunk(source) | ||||
if not c: | ||||
break | ||||
Augie Fackler
|
r10430 | elif progress is not None: | ||
progress() | ||||
Thomas Arendsen Hein
|
r1981 | yield c | ||
Matt Mackall
|
r5368 | def chunkheader(length): | ||
Greg Ward
|
r9437 | """return a changegroup chunk header (string)""" | ||
Matt Mackall
|
r5368 | return struct.pack(">l", length + 4) | ||
Thomas Arendsen Hein
|
r1981 | |||
def closechunk(): | ||||
Greg Ward
|
r9437 | """return a changegroup chunk header (string) for a zero-length chunk""" | ||
Thomas Arendsen Hein
|
r1981 | return struct.pack(">l", 0) | ||
Matt Mackall
|
r3659 | class nocompress(object): | ||
def compress(self, x): | ||||
return x | ||||
def flush(self): | ||||
return "" | ||||
Matt Mackall
|
r3662 | bundletypes = { | ||
Benoit Boissinot
|
r3704 | "": ("", nocompress), | ||
"HG10UN": ("HG10UN", nocompress), | ||||
Alexis S. L. Carvalho
|
r3762 | "HG10BZ": ("HG10", lambda: bz2.BZ2Compressor()), | ||
"HG10GZ": ("HG10GZ", lambda: zlib.compressobj()), | ||||
Matt Mackall
|
r3662 | } | ||
Dirkjan Ochtman
|
r10356 | def collector(cl, mmfs, files): | ||
# Gather information about changeset nodes going out in a bundle. | ||||
# We want to gather manifests needed and filelogs affected. | ||||
def collect(node): | ||||
c = cl.read(node) | ||||
Benoit Boissinot
|
r11648 | files.update(c[3]) | ||
Dirkjan Ochtman
|
r10356 | mmfs.setdefault(c[0], node) | ||
return collect | ||||
Martin Geisler
|
r9087 | # hgweb uses this list to communicate its preferred type | ||
Dirkjan Ochtman
|
r6152 | bundlepriority = ['HG10GZ', 'HG10BZ', 'HG10UN'] | ||
Thomas Arendsen Hein
|
r3706 | def writebundle(cg, filename, bundletype): | ||
Matt Mackall
|
r3659 | """Write a bundle file and return its filename. | ||
Existing files will not be overwritten. | ||||
If no filename is specified, a temporary file is created. | ||||
bz2 compression can be turned off. | ||||
The bundle file will be deleted in case of errors. | ||||
""" | ||||
fh = None | ||||
cleanup = None | ||||
try: | ||||
if filename: | ||||
fh = open(filename, "wb") | ||||
else: | ||||
fd, filename = tempfile.mkstemp(prefix="hg-bundle-", suffix=".hg") | ||||
fh = os.fdopen(fd, "wb") | ||||
cleanup = filename | ||||
Thomas Arendsen Hein
|
r3706 | header, compressor = bundletypes[bundletype] | ||
Benoit Boissinot
|
r3704 | fh.write(header) | ||
z = compressor() | ||||
Matt Mackall
|
r3662 | |||
Matt Mackall
|
r3659 | # parse the changegroup data, otherwise we will block | ||
# in case of sshrepo because we don't know the end of the stream | ||||
# an empty chunkiter is the end of the changegroup | ||||
Alexis S. L. Carvalho
|
r5906 | # a changegroup has at least 2 chunkiters (changelog and manifest). | ||
# after that, an empty chunkiter is the end of the changegroup | ||||
Matt Mackall
|
r3659 | empty = False | ||
Alexis S. L. Carvalho
|
r5906 | count = 0 | ||
while not empty or count <= 2: | ||||
Matt Mackall
|
r3659 | empty = True | ||
Alexis S. L. Carvalho
|
r5906 | count += 1 | ||
Matt Mackall
|
r3659 | for chunk in chunkiter(cg): | ||
empty = False | ||||
Matt Mackall
|
r5368 | fh.write(z.compress(chunkheader(len(chunk)))) | ||
pos = 0 | ||||
while pos < len(chunk): | ||||
next = pos + 2**20 | ||||
fh.write(z.compress(chunk[pos:next])) | ||||
pos = next | ||||
Matt Mackall
|
r3659 | fh.write(z.compress(closechunk())) | ||
fh.write(z.flush()) | ||||
cleanup = None | ||||
return filename | ||||
finally: | ||||
if fh is not None: | ||||
fh.close() | ||||
if cleanup is not None: | ||||
os.unlink(cleanup) | ||||
Matt Mackall
|
r3660 | |||
Dirkjan Ochtman
|
r6154 | def unbundle(header, fh): | ||
if header == 'HG10UN': | ||||
return fh | ||||
elif not header.startswith('HG'): | ||||
# old client with uncompressed bundle | ||||
def generator(f): | ||||
yield header | ||||
for chunk in f: | ||||
yield chunk | ||||
elif header == 'HG10GZ': | ||||
def generator(f): | ||||
zd = zlib.decompressobj() | ||||
for chunk in f: | ||||
yield zd.decompress(chunk) | ||||
elif header == 'HG10BZ': | ||||
Matt Mackall
|
r3660 | def generator(f): | ||
zd = bz2.BZ2Decompressor() | ||||
zd.decompress("BZ") | ||||
for chunk in util.filechunkiter(f, 4096): | ||||
yield zd.decompress(chunk) | ||||
Dirkjan Ochtman
|
r6154 | return util.chunkbuffer(generator(fh)) | ||
Matt Mackall
|
r3660 | |||
Dirkjan Ochtman
|
r6154 | def readbundle(fh, fname): | ||
header = fh.read(6) | ||||
if not header.startswith('HG'): | ||||
raise util.Abort(_('%s: not a Mercurial bundle file') % fname) | ||||
if not header.startswith('HG10'): | ||||
raise util.Abort(_('%s: unknown bundle version') % fname) | ||||
elif header not in bundletypes: | ||||
raise util.Abort(_('%s: unknown bundle compression type') % fname) | ||||
return unbundle(header, fh) | ||||