trex5790 commited on
Commit
00732ed
·
verified ·
1 Parent(s): f75ca5a

Update config.json

Browse files

replace "bnb_4bit_quant_type": "nf4", with "bnb_4bit_quant_type": "gptq",

Files changed (1) hide show
  1. config.json +1 -2
config.json CHANGED
@@ -19,7 +19,7 @@
19
  "pretraining_tp": 1,
20
  "quantization_config": {
21
  "bnb_4bit_compute_dtype": "float16",
22
- "bnb_4bit_quant_type": "nf4",
23
  "bnb_4bit_use_double_quant": true,
24
  "llm_int8_enable_fp32_cpu_offload": false,
25
  "llm_int8_has_fp16_weight": false,
@@ -27,7 +27,6 @@
27
  "llm_int8_threshold": 6.0,
28
  "load_in_4bit": true,
29
  "load_in_8bit": false,
30
- "quantization": "gptq",
31
  "quant_method": "bitsandbytes"
32
  },
33
  "rms_norm_eps": 1e-05,
 
19
  "pretraining_tp": 1,
20
  "quantization_config": {
21
  "bnb_4bit_compute_dtype": "float16",
22
+ "bnb_4bit_quant_type": "gptq",
23
  "bnb_4bit_use_double_quant": true,
24
  "llm_int8_enable_fp32_cpu_offload": false,
25
  "llm_int8_has_fp16_weight": false,
 
27
  "llm_int8_threshold": 6.0,
28
  "load_in_4bit": true,
29
  "load_in_8bit": false,
 
30
  "quant_method": "bitsandbytes"
31
  },
32
  "rms_norm_eps": 1e-05,