Download quantization_config.json from Vishva007/clef-flash-W4A16-AutoRound: direct link, hf CLI and curl.
- Browser
- Download file 411 Bytes
-
https://huggingface.co/Vishva007/clef-flash-W4A16-AutoRound/resolve/main/quantization_config.json
- Command line
-
hf download hf://Vishva007/clef-flash-W4A16-AutoRound/quantization_config.json
-
curl -L -o quantization_config.json https://huggingface.co/Vishva007/clef-flash-W4A16-AutoRound/resolve/main/quantization_config.json
411 Bytes
| { | |
| "bits": 4, | |
| "data_type": "int", | |
| "group_size": 32, | |
| "sym": true, | |
| "iters": 800, | |
| "autoround_version": "0.16.0", | |
| "block_name_to_quantize": "model.language_model.layers", | |
| "quant_method": "auto-round", | |
| "packing_format": "auto_round:auto_gptq", | |
| "extra_config": { | |
| ".*\\*\\.linear_conv.*": { | |
| "data_type": "bfloat16" | |
| }, | |
| ".*\\*\\.conv1d.*": { | |
| "data_type": "bfloat16" | |
| } | |
| } | |
| } |