Fix GPTQ converter (#423)

* Fix GPTQ converter * Fix comment --------- Co-authored-by: Georgi Gerganov <ggerganov@gmail.com>
2024-12-25 02:44:36 +00:00 · 2023-03-23 10:18:13 -10:00 · 2023-03-23 10:18:13 -10:00 · 20a1a4e09c
commit 20a1a4e09c
parent ad072fc5ad
1 changed files with 8 additions and 13 deletions
--- a/convert-gptq-to-ggml.py
+++ b/convert-gptq-to-ggml.py
@ -36,7 +36,8 @@ fname_out = sys.argv[3]
 fout = open(fname_out, "wb")
-fout.write(struct.pack("i", 0x67676d6c)) # magic: ggml in hex
+fout.write(struct.pack("i", 0x67676d66)) # magic: ggmf in hex
 fout.write(struct.pack("i", 1)) # file version
 fout.write(struct.pack("i", n_vocab))
 fout.write(struct.pack("i", n_embd))
 fout.write(struct.pack("i", n_mult))
@ -49,27 +50,21 @@ fout.write(struct.pack("i", 4))
 # This loop unchanged from convert-pth-to-ggml.py:
 for i in range(tokenizer.vocab_size()):
    if tokenizer.is_unknown(i):
        # "<unk>" token (translated as ??)
        text = " \u2047 ".encode("utf-8")
        fout.write(struct.pack("i", len(text)))
        fout.write(text)
    elif tokenizer.is_control(i):
-        # "<s>"/"</s>" tokens
+        text = b""
        fout.write(struct.pack("i", 0))
    elif tokenizer.is_byte(i):
        # "<U+XX>" tokens (which may be invalid UTF-8)
        piece = tokenizer.id_to_piece(i)
        if len(piece) != 6:
-            print("Invalid token: " + piece)
+            print(f"Invalid token: {piece}")
            sys.exit(1)
        byte_value = int(piece[3:-1], 16)
-        fout.write(struct.pack("i", 1))
+        text = struct.pack("B", byte_value)
        fout.write(struct.pack("B", byte_value))
    else:
        # normal token. Uses U+2581 (LOWER ONE EIGHTH BLOCK) to represent spaces.
        text = tokenizer.id_to_piece(i).replace("\u2581", " ").encode("utf-8")
-        fout.write(struct.pack("i", len(text)))
+    fout.write(struct.pack("i", len(text)))
-        fout.write(text)
+    fout.write(text)
    fout.write(struct.pack("f", tokenizer.get_score(i)))
 def write_header(shape, dst_name, ftype_cur):
    sname = dst_name.encode('utf-8')