{ "bits": 4, "data_type": "int", "group_size": 32, "sym": true, "iters": 1000, "static_attention_granularity": "tensor", "static_kv_granularity": "tensor", "autoround_version": "0.16.0", "dynamic": { "+:.*mtp.*": { "bits": 4, "group_size": 32, "sym": true }, "+:.*mtp\\.fc.*": { "bits": 4, "group_size": 32, "sym": true } }, "lm_head": false, "provider": "auto-round", "quant_method": "gptq", "desc_act": false, "true_sequential": false, "damp_percent": 0.01 }