convert : fix autoawq gemma (#6704)

author Zheng.Deng <redacted>

Tue, 16 Apr 2024 20:51:07 +0000 (04:51 +0800)

committer GitHub <redacted>

Tue, 16 Apr 2024 20:51:07 +0000 (23:51 +0300)
author Zheng.Deng <redacted>
Tue, 16 Apr 2024 20:51:07 +0000 (04:51 +0800)
committer GitHub <redacted>
Tue, 16 Apr 2024 20:51:07 +0000 (23:51 +0300)
diff --git a/convert-hf-to-gguf.py b/convert-hf-to-gguf.py

index f321d77de11f87ca6296d1a7accdac7034542271..c14186abbc2a6cb6df585e6cf1bd2c195a960397 100755 (executable)
--- a/convert-hf-to-gguf.py
+++ b/convert-hf-to-gguf.py
@@ -2458,6 +2458,12 @@ class GemmaModel(Model):
          tensor_map = gguf.get_tensor_name_map(self.model_arch, block_count)
  
          for name, data_torch in self.get_tensors():
+            # lm_head is not used in llama.cpp, while autoawq will include this tensor in model
+            # To prevent errors, skip loading lm_head.weight.
+            if name == "lm_head.weight":
+                print(f"Skipping get tensor {name!r} in safetensors so that convert can end normally.")
+                continue
+
              old_dtype = data_torch.dtype
  
              # convert any unsupported data types to float32
author	Zheng.Deng <redacted>
	Tue, 16 Apr 2024 20:51:07 +0000 (04:51 +0800)
committer	GitHub <redacted>
	Tue, 16 Apr 2024 20:51:07 +0000 (23:51 +0300)