jkim96 commited on
Commit
bd8ca46
·
verified ·
1 Parent(s): 5b3f20b

Upload dashq quantized checkpoint (INT3, g64, scale_zero_dtype=float16, zero-shot avg=70.6)

Browse files
README.md CHANGED
@@ -38,6 +38,7 @@ model, tokenizer = load_quantized(
38
  | Base model | `Qwen/Qwen3.6-27B` |
39
  | Bits | `3` |
40
  | Group size | `64` |
 
41
  | Calibration dataset | `wikitext2` |
42
  | Calibration samples | `128` |
43
  | Sequence length | `2048` |
@@ -48,14 +49,15 @@ model, tokenizer = load_quantized(
48
 
49
  | Metric | Value |
50
  | --- | ---: |
51
- | `wikitext2_ppl` | 7.7746 |
52
- | `arc_challenge` | 60.7509 |
53
- | `arc_easy` | 79.2929 |
54
- | `commonsense_qa` | 63.1450 |
55
- | `hellaswag` | 82.3442 |
56
- | `lambada_openai` | 74.6167 |
57
- | `openbookqa` | 44.8000 |
58
- | `piqa` | 82.0457 |
59
- | `truthfulqa_mc2` | 54.7740 |
 
60
  | `winogrande` | 76.9534 |
61
 
 
38
  | Base model | `Qwen/Qwen3.6-27B` |
39
  | Bits | `3` |
40
  | Group size | `64` |
41
+ | Scale/zero dtype | `float16` |
42
  | Calibration dataset | `wikitext2` |
43
  | Calibration samples | `128` |
44
  | Sequence length | `2048` |
 
49
 
50
  | Metric | Value |
51
  | --- | ---: |
52
+ | `wikitext2_ppl` | 7.7131 |
53
+ | `zero-shot accuracy avg` | 70.5799 |
54
+ | `arc_challenge` | 60.1536 |
55
+ | `arc_easy` | 78.4512 |
56
+ | `commonsense_qa` | 79.4431 |
57
+ | `hellaswag` | 81.8861 |
58
+ | `lambada_openai` | 75.1213 |
59
+ | `openbookqa` | 45.4000 |
60
+ | `piqa` | 81.8825 |
61
+ | `truthfulqa_mc2` | 55.9281 |
62
  | `winogrande` | 76.9534 |
63
 
dashq_config.json CHANGED
@@ -10,6 +10,7 @@
10
  "low_memory_optimization": false,
11
  "moe_hessian_scope": "shared",
12
  "n_samples": 128,
 
13
  "symmetric": false,
14
  "use_error_compensation": true,
15
  "use_optimal_shrinkage": true,
@@ -5974,19 +5975,18 @@
5974
  "Model": "Qwen/Qwen3.6-27B",
5975
  "ModelSizeGB": 17.27475596,
5976
  "OriginalSizeGB": 55.5630064,
5977
- "PPL": 7.774606227874756,
5978
- "Params": "{'bits': 3, 'group_size': 64, 'n_samples': 128, 'moe_hessian_scope': 'shared', 'use_error_compensation': True, 'use_optimal_shrinkage': True, 'use_weighted_quantization': True, 'symmetric': False, 'low_memory_optimization': False}",
5979
- "QuantTime": 884.1916358470917,
5980
- "arc_challenge": 60.75085324232082,
5981
- "arc_easy": 79.29292929292929,
5982
- "boolq": 67.55351681957187,
5983
- "commonsense_qa": 63.14496314496314,
5984
- "hellaswag": 82.34415455088627,
5985
- "lambada_openai": 74.61672811954202,
5986
- "openbookqa": 44.800000000000004,
5987
- "piqa": 82.04570184983679,
5988
- "social_iqa": 68.3204134366925,
5989
- "truthfulqa_mc2": 54.774021091223176,
5990
- "winogrande": 76.95343330702447
5991
  }
5992
  }
 
10
  "low_memory_optimization": false,
11
  "moe_hessian_scope": "shared",
12
  "n_samples": 128,
13
+ "scale_zero_dtype": "float16",
14
  "symmetric": false,
15
  "use_error_compensation": true,
16
  "use_optimal_shrinkage": true,
 
5975
  "Model": "Qwen/Qwen3.6-27B",
5976
  "ModelSizeGB": 17.27475596,
5977
  "OriginalSizeGB": 55.5630064,
5978
+ "PPL": 7.71309757232666,
5979
+ "Params": "{'bits': 3, 'group_size': 64, 'scale_zero_dtype': 'float16', 'n_samples': 128, 'moe_hessian_scope': 'shared', 'use_error_compensation': True, 'use_optimal_shrinkage': True, 'use_weighted_quantization': True, 'symmetric': False, 'low_memory_optimization': False}",
5980
+ "QuantTime": 894.5826256275177,
5981
+ "arc_challenge": 60.153583617747444,
5982
+ "arc_easy": 78.45117845117845,
5983
+ "commonsense_qa": 79.44307944307944,
5984
+ "hellaswag": 81.88607847042422,
5985
+ "lambada_openai": 75.12128856976518,
5986
+ "openbookqa": 45.4,
5987
+ "piqa": 81.88248095756256,
5988
+ "truthfulqa_mc2": 55.9280749996677,
5989
+ "winogrande": 76.95343330702447,
5990
+ "zeroshot_avg": 70.57991086849438
 
5991
  }
5992
  }
model-00002-of-00004.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:6bb4da5b9691b1b833ed3c930800378c7ad5315b367b287951099cb3120969c3
3
  size 4968942552
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:69e03e3443d8fe5513a46fb33c09b94c8ec6122ec0d06534756a879471533a67
3
  size 4968942552
model-00003-of-00004.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:aa75309016b9f6260412b783767e747d9286f496a952b64c1542801c7522a7ad
3
  size 4995401320
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:cfb922d084d81c3cbecdc222c7a1112957077d6f88b09ae5978e87d0d2f71bb5
3
  size 4995401320
model-00004-of-00004.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:4e27298c29c957fab5490ce7019c8d54f656797a66962f64492efa7fd4a8b1e5
3
  size 4767615160
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:52686153606d162c95b2276e99ba1532b6202e92687ef264f748700f773cb281
3
  size 4767615160