fxmarty-amd commited on
Commit
1dfea05
·
verified ·
1 Parent(s): df71383

Upload config.json with huggingface_hub

Browse files
Files changed (1) hide show
  1. config.json +5 -3
config.json CHANGED
@@ -38,10 +38,12 @@
38
  },
39
  "format": "pack-quantized",
40
  "ignore": [
41
- "lm_head",
42
  "re:.*self_attn.*",
43
  "re:.*shared_experts.*",
44
- "re:.*mlp\\.(gate|up|gate_up|down)_proj.*"
 
 
 
45
  ],
46
  "kv_cache_scheme": null,
47
  "quant_method": "compressed-tensors",
@@ -222,4 +224,4 @@
222
  "vt_num_attention_heads": 16,
223
  "vt_num_hidden_layers": 27
224
  }
225
- }
 
38
  },
39
  "format": "pack-quantized",
40
  "ignore": [
 
41
  "re:.*self_attn.*",
42
  "re:.*shared_experts.*",
43
+ "re:.*mlp\\.(gate|up|gate_up|down)_proj.*",
44
+ "re:.*lm_head.*",
45
+ "re:vision_tower.*",
46
+ "re:mm_projector.*"
47
  ],
48
  "kv_cache_scheme": null,
49
  "quant_method": "compressed-tensors",
 
224
  "vt_num_attention_heads": 16,
225
  "vt_num_hidden_layers": 27
226
  }
227
+ }