/brz/remove-bazaar : revision 0.200.1036

To get this branch, use:

bzr branch
http://gegoxaren.bato24.eu/bzr/brz/remove-bazaar

« back to all changes in this revision

Viewing changes to bzrlib/cache.py

Committer: Jelmer Vernooij
Date: 2010-09-11 17:48:45 UTC
mto: (0.312.1 master) (6883.23.1 bundle-git)
mto: This revision was merged to the branch mainline in revision 6960.
Revision ID: jelmer@samba.org-20100911174845-ro06bb6gg6ws8jit

More work on roundtrip push support.

files added:
.bzrignore

COPYING

HACKING

INSTALL

Makefile

NEWS

README

TODO

__init__.py

branch.py

bzr-receive-pack

bzr-upload-pack

cache.py

commands.py

commit.py

config.py

dir.py

errors.py

fetch.py

help.py

hg.py

info.py

inventory.py

mapping.py

notes

notes/git-serve.txt

notes/mapping.txt

notes/roundtripping.txt

object_store.py

push.py

refs.py

remote.py

repository.py

revspec.py

roundtrip.py

send.py

server.py

setup.py

tests

tests/__init__.py

tests/test_blackbox.py

tests/test_branch.py

tests/test_builder.py

tests/test_cache.py

tests/test_dir.py

tests/test_fetch.py

tests/test_mapping.py

tests/test_object_store.py

tests/test_push.py

tests/test_refs.py

tests/test_remote.py

tests/test_repository.py

tests/test_revspec.py

tests/test_roundtrip.py

tests/test_transportgit.py

transportgit.py

tree.py

versionedfiles.py

workingtree.py

files removed:
.bzrignore

.rsyncexclude

NEWS

README

TODO

build-api

bzrlib

bzrlib/__init__.py

bzrlib/add.py

bzrlib/branch.py

bzrlib/cache.py

bzrlib/check.py

bzrlib/commands.py

bzrlib/diff.py

bzrlib/errors.py

bzrlib/help.py

bzrlib/info.py

bzrlib/inventory.py

bzrlib/log.py

bzrlib/mdiff.py

bzrlib/newinventory.py

bzrlib/osutils.py

bzrlib/remotebranch.py

bzrlib/revfile.py

bzrlib/revision.py

bzrlib/status.py

bzrlib/store.py

bzrlib/tests.py

bzrlib/textinv.py

bzrlib/textui.py

bzrlib/trace.py

bzrlib/tree.py

bzrlib/xml.py

contrib

contrib/add-bzr-to-baz

contrib/bash

contrib/bash/bzr

contrib/zsh

contrib/zsh/_bzr

doc/Makefile

doc/adoption.txt

doc/bitkeeper.txt

doc/changelogs.txt

doc/cherry-picking.txt

doc/cmdref.txt

doc/common-format.txt

doc/compared-aegis.txt

doc/compared-codeville.txt

doc/compared-cvsnt.txt

doc/compared-opencm.txt

doc/compared-prcs.txt

doc/compared-teamware.txt

doc/compression.txt

doc/config-specs.txt

doc/conflicts.txt

doc/costs.txt

doc/darcs.txt

doc/deadly-sins.txt

doc/default.css

doc/design.txt

doc/extra-commands.txt

doc/formats.txt

doc/hashes.txt

doc/ignore.txt

doc/index.txt

doc/interrupted.txt

doc/intro.txt

doc/inventory.txt

doc/join-branches.txt

doc/kill-version.txt

doc/layers.txt

doc/library-interface.txt

doc/merge.txt

doc/mirroring.txt

doc/monotone.txt

doc/news.txt

doc/optional-edit.txt

doc/partial-commit.txt

doc/pool.txt

doc/purpose.txt

doc/python.txt

doc/quilt.txt

doc/quotes.txt

doc/random.txt

doc/requirements.txt

doc/revfile.txt

doc/revision-syntax.txt

doc/rollup.txt

doc/scalability.txt

doc/security.txt

doc/shared-branches.txt

doc/short-demo.txt

doc/supportability.txt

doc/svk.txt

doc/switch-in-branch.txt

doc/tagging.txt

doc/taxonomy.txt

doc/thanks.txt

doc/todo-from-arch.txt

doc/unchanged.txt

doc/unrelated-merge.txt

doc/usability.txt

doc/use-cases.txt

doc/web-interface.txt

doc/workflow.txt

doc/yaml.txt

elementtree

elementtree/ElementTree.py

elementtree/__init__.py

notes

notes/new-inventory-sample.xml

notes/performance.txt

setup.py

testbzr

urlgrabber

urlgrabber/__init__.py

urlgrabber/byterange.py

urlgrabber/grabber.py

urlgrabber/keepalive.py

urlgrabber/mirror.py

urlgrabber/progress.py

Show diffs side-by-side

added added

removed removed

bzrlib/cache.py

# This program is free software; you can redistribute it and/or modify

# it under the terms of the GNU General Public License as published by

# the Free Software Foundation; either version 2 of the License, or

# (at your option) any later version.

# This program is distributed in the hope that it will be useful,

# but WITHOUT ANY WARRANTY; without even the implied warranty of

# MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the

# GNU General Public License for more details.

# You should have received a copy of the GNU General Public License

# along with this program; if not, write to the Free Software

# Foundation, Inc., 59 Temple Place, Suite 330, Boston, MA 02111-1307 USA

import stat, os, sha, time

from binascii import b2a_qp, a2b_qp

from trace import mutter

# file fingerprints are: (path, size, mtime, ctime, ino, dev).

# if this is the same for this file as in the previous revision, we

# assume the content is the same and the SHA-1 is the same.

# This is stored in a fingerprint file that also contains the file-id

# and the content SHA-1.

# Thus for any given file we can quickly get the SHA-1, either from

# the cache or if the cache is out of date.

# At the moment this is stored in a simple textfile; it might be nice

# to use a tdb instead.

# What we need:

# build a new cache from scratch

# load cache, incrementally update it

# TODO: Have a paranoid mode where we always compare the texts and

# always recalculate the digest, to trap modification without stat

# change and SHA collisions.

def fingerprint(path, abspath):

try:

fs = os.lstat(abspath)

except OSError:

# might be missing, etc

return None

if stat.S_ISDIR(fs.st_mode):

return None

return (fs.st_size, fs.st_mtime,

fs.st_ctime, fs.st_ino, fs.st_dev)

def write_cache(branch, entry_iter):

outf = branch.controlfile('work-cache.tmp', 'wt')

for entry in entry_iter:

outf.write(entry[0] + ' ' + entry[1] + ' ')

outf.write(b2a_qp(entry[2], True))

outf.write(' %d %d %d %d %d\n' % entry[3:])

outf.close()

os.rename(branch.controlfilename('work-cache.tmp'),

branch.controlfilename('work-cache'))

def load_cache(branch):

cache = {}

try:

cachefile = branch.controlfile('work-cache', 'rt')

except IOError:

return cache

for l in cachefile:

f = l.split(' ')

file_id = f[0]

if file_id in cache:

raise BzrError("duplicated file_id in cache: {%s}" % file_id)

cache[file_id] = (f[0], f[1], a2b_qp(f[2])) + tuple([long(x) for x in f[3:]])

return cache

def _files_from_inventory(inv):

for path, ie in inv.iter_entries():

if ie.kind != 'file':

continue

yield ie.file_id, path

100

101

102

def build_cache(branch):

103

inv = branch.read_working_inventory()

104

105

cache = {}

106

_update_cache_from_list(branch, cache, _files_from_inventory(inv))

107

108

109

110

def update_cache(branch, inv):

111

# TODO: It's supposed to be faster to stat the files in order by inum.

112

# We don't directly know the inum of the files of course but we do

113

# know where they were last sighted, so we can sort by that.

114

115

cache = load_cache(branch)

116

return _update_cache_from_list(branch, cache, _files_from_inventory(inv))

117

118

119

120

def _update_cache_from_list(branch, cache, to_update):

121

"""Update the cache to have info on the named files.

122

123

to_update is a sequence of (file_id, path) pairs.

124

"""

125

hardcheck = dirty = 0

126

for file_id, path in to_update:

127

fap = branch.abspath(path)

128

fp = fingerprint(fap, path)

129

cacheentry = cache.get(file_id)

130

131

if fp == None: # not here

132

if cacheentry:

133

del cache[file_id]

134

dirty += 1

135

continue

136

137

if cacheentry and (cacheentry[3:] == fp):

138

continue # all stat fields unchanged

139

140

hardcheck += 1

141

142

dig = sha.new(file(fap, 'rb').read()).hexdigest()

143

144

if cacheentry == None or dig != cacheentry[1]:

145

# if there was no previous entry for this file, or if the

146

# SHA has changed, then update the cache

147

cacheentry = (file_id, dig, path) + fp

148

cache[file_id] = cacheentry

149

dirty += 1

150

151

mutter('work cache: read %d files, %d changed' % (hardcheck, dirty))

152

153

if dirty:

154

write_cache(branch, cache.itervalues())

155

156

return cache

Older »