2
This is just a utility that I use to extract the projector for quantized models.
3
It is NOT necessary at all to train, or run inference/serve demos.
4
Use this script ONLY if you fully understand its implications.
12
from collections import defaultdict
16
parser = argparse.ArgumentParser(description='Extract MMProjector weights')
17
parser.add_argument('--model-path', type=str, help='model folder')
18
parser.add_argument('--output', type=str, help='output file')
19
args = parser.parse_args()
23
if __name__ == '__main__':
26
keys_to_match = ['mm_projector']
27
ckpt_to_key = defaultdict(list)
29
model_indices = json.load(open(os.path.join(args.model_path, 'pytorch_model.bin.index.json')))
30
for k, v in model_indices['weight_map'].items():
31
if any(key_match in k for key_match in keys_to_match):
32
ckpt_to_key[v].append(k)
33
except FileNotFoundError:
35
v = 'pytorch_model.bin'
36
for k in torch.load(os.path.join(args.model_path, v), map_location='cpu').keys():
37
if any(key_match in k for key_match in keys_to_match):
38
ckpt_to_key[v].append(k)
42
for ckpt_name, weight_keys in ckpt_to_key.items():
43
ckpt = torch.load(os.path.join(args.model_path, ckpt_name), map_location='cpu')
45
loaded_weights[k] = ckpt[k]
47
torch.save(loaded_weights, args.output)