jkim96 commited on
Commit
98efcc6
·
verified ·
1 Parent(s): 7faf2fb

Upload dashq quantized checkpoint (INT4, g128, scale_zero_dtype=float16, zero-shot avg=71.2)

Browse files
README.md CHANGED
@@ -38,6 +38,7 @@ model, tokenizer = load_quantized(
38
  | Base model | `Qwen/Qwen3.6-27B` |
39
  | Bits | `4` |
40
  | Group size | `128` |
 
41
  | Calibration dataset | `wikitext2` |
42
  | Calibration samples | `128` |
43
  | Sequence length | `2048` |
@@ -48,13 +49,15 @@ model, tokenizer = load_quantized(
48
 
49
  | Metric | Value |
50
  | --- | ---: |
51
- | `wikitext2_ppl` | 7.5355 |
52
- | `arc_challenge` | 59.9829 |
53
- | `arc_easy` | 76.3889 |
54
- | `commonsense_qa` | 85.9951 |
55
- | `hellaswag` | 83.3201 |
56
- | `lambada_openai` | 74.3450 |
57
- | `openbookqa` | 45.0000 |
58
- | `piqa` | 82.6986 |
59
- | `truthfulqa_mc2` | 56.6694 |
60
- | `winogrande` | 77.7427 |
 
 
 
38
  | Base model | `Qwen/Qwen3.6-27B` |
39
  | Bits | `4` |
40
  | Group size | `128` |
41
+ | Scale/zero dtype | `float16` |
42
  | Calibration dataset | `wikitext2` |
43
  | Calibration samples | `128` |
44
  | Sequence length | `2048` |
 
49
 
50
  | Metric | Value |
51
  | --- | ---: |
52
+ | `wikitext2_ppl` | 7.4794 |
53
+ | `zero-shot accuracy avg` | 71.2425 |
54
+ | `arc_challenge` | 59.5563 |
55
+ | `arc_easy` | 76.0101 |
56
+ | `commonsense_qa` | 85.8313 |
57
+ | `hellaswag` | 83.4396 |
58
+ | `lambada_openai` | 74.7720 |
59
+ | `openbookqa` | 44.2000 |
60
+ | `piqa` | 82.8618 |
61
+ | `truthfulqa_mc2` | 56.6108 |
62
+ | `winogrande` | 77.9006 |
63
+
dashq_config.json CHANGED
@@ -10,6 +10,7 @@
10
  "low_memory_optimization": false,
11
  "moe_hessian_scope": "shared",
12
  "n_samples": 128,
 
13
  "symmetric": false,
14
  "use_error_compensation": true,
15
  "use_optimal_shrinkage": true,
@@ -5974,20 +5975,18 @@
5974
  "Model": "Qwen/Qwen3.6-27B",
5975
  "ModelSizeGB": 18.948856152,
5976
  "OriginalSizeGB": 55.5630064,
5977
- "PPL": 7.53549337387085,
5978
- "Params": "{'bits': 4, 'group_size': 128, 'n_samples': 128, 'moe_hessian_scope': 'shared', 'use_error_compensation': True, 'use_optimal_shrinkage': True, 'use_weighted_quantization': True, 'symmetric': False, 'low_memory_optimization': False}",
5979
- "QuantTime": 898.2053005695343,
5980
- "arc_challenge": 59.98293515358362,
5981
- "arc_easy": 76.38888888888889,
5982
- "boolq": 77.2782874617737,
5983
- "commonsense_qa": 85.995085995086,
5984
- "hellaswag": 83.3200557657837,
5985
- "lambada_openai": 74.345041723268,
5986
- "openbookqa": 45.0,
5987
- "piqa": 82.69858541893362,
5988
- "social_iqa": 65.83979328165374,
5989
- "truthfulqa_mc2": 56.669406789395374,
5990
- "winogrande": 77.7426992896606,
5991
- "zeroshot_avg": 71.38734361527521
5992
  }
5993
- }
 
10
  "low_memory_optimization": false,
11
  "moe_hessian_scope": "shared",
12
  "n_samples": 128,
13
+ "scale_zero_dtype": "float16",
14
  "symmetric": false,
15
  "use_error_compensation": true,
16
  "use_optimal_shrinkage": true,
 
5975
  "Model": "Qwen/Qwen3.6-27B",
5976
  "ModelSizeGB": 18.948856152,
5977
  "OriginalSizeGB": 55.5630064,
5978
+ "PPL": 7.479367256164551,
5979
+ "Params": "{'bits': 4, 'group_size': 128, 'scale_zero_dtype': 'float16', 'n_samples': 128, 'moe_hessian_scope': 'shared', 'use_error_compensation': True, 'use_optimal_shrinkage': True, 'use_weighted_quantization': True, 'symmetric': False, 'low_memory_optimization': False}",
5980
+ "QuantTime": 839.7647020816803,
5981
+ "arc_challenge": 59.55631399317406,
5982
+ "arc_easy": 76.01010101010101,
5983
+ "commonsense_qa": 85.83128583128583,
5984
+ "hellaswag": 83.43955387373033,
5985
+ "lambada_openai": 74.77197748884146,
5986
+ "openbookqa": 44.2,
5987
+ "piqa": 82.86180631120783,
5988
+ "truthfulqa_mc2": 56.610824269727146,
5989
+ "winogrande": 77.90055248618785,
5990
+ "zeroshot_avg": 71.24249058491728
 
 
5991
  }
5992
+ }
model-00002-of-00005.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:6bd1ef08a93ec894b46d7c7e3c84de93ee51627f49490319104054777dc23d02
3
  size 4997583600
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:72cbb0bfce29be4bef20951afea7ffbbd29f9fd4e38b5c661c623a6db7c46a5a
3
  size 4997583600
model-00003-of-00005.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:385a8da079a4c98292d0527a1165d170f17f6fe3bf27153e930611e6e8f149ca
3
  size 4981282880
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:9d0eb5a07bd954735323559955be5c33c1ea24a234491b8de7f1add388c7ad30
3
  size 4981282880
model-00004-of-00005.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:01e18b59d224c133b4bf441394ae629c7c72d62ad58ce4bc2837698742aa2179
3
  size 4962151088
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:275636a91121fdbda2cd480cfb720066ac2b79a390a18a6ee69eaeea77f1b263
3
  size 4962151088
model-00005-of-00005.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:f9042bbfb1dae6c639cc37e18bb1e70f99979379b0dde16b11507a7ade370512
3
  size 1465041656
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:261d3088c8eb0fb6ed386b3db63447589a281f6b55cf5da72884b6ce45e68981
3
  size 1465041656