jkim96 commited on
Commit
de8bfc9
·
verified ·
1 Parent(s): 7d666af

Upload dashq quantized checkpoint (INT3, g128, scale_zero_dtype=float16, zero-shot avg=70.8)

Browse files
README.md CHANGED
@@ -38,6 +38,7 @@ model, tokenizer = load_quantized(
38
  | Base model | `Qwen/Qwen3.6-27B` |
39
  | Bits | `3` |
40
  | Group size | `128` |
 
41
  | Calibration dataset | `wikitext2` |
42
  | Calibration samples | `128` |
43
  | Sequence length | `2048` |
@@ -48,13 +49,15 @@ model, tokenizer = load_quantized(
48
 
49
  | Metric | Value |
50
  | --- | ---: |
51
- | `wikitext2_ppl` | 7.7625 |
52
- | `arc_challenge` | 59.8123 |
53
- | `arc_easy` | 77.1044 |
54
- | `commonsense_qa` | 84.8485 |
55
- | `hellaswag` | 82.2645 |
56
- | `lambada_openai` | 75.6647 |
57
- | `openbookqa` | 45.4000 |
58
- | `piqa` | 81.9913 |
59
- | `truthfulqa_mc2` | 54.7219 |
60
- | `winogrande` | 76.7167 |
 
 
 
38
  | Base model | `Qwen/Qwen3.6-27B` |
39
  | Bits | `3` |
40
  | Group size | `128` |
41
+ | Scale/zero dtype | `float16` |
42
  | Calibration dataset | `wikitext2` |
43
  | Calibration samples | `128` |
44
  | Sequence length | `2048` |
 
49
 
50
  | Metric | Value |
51
  | --- | ---: |
52
+ | `wikitext2_ppl` | 7.7554 |
53
+ | `zero-shot accuracy avg` | 70.7634 |
54
+ | `arc_challenge` | 58.7884 |
55
+ | `arc_easy` | 75.5471 |
56
+ | `commonsense_qa` | 86.1589 |
57
+ | `hellaswag` | 82.5931 |
58
+ | `lambada_openai` | 74.7720 |
59
+ | `openbookqa` | 44.8000 |
60
+ | `piqa` | 82.4266 |
61
+ | `truthfulqa_mc2` | 54.9096 |
62
+ | `winogrande` | 76.8745 |
63
+
dashq_config.json CHANGED
@@ -10,6 +10,7 @@
10
  "low_memory_optimization": false,
11
  "moe_hessian_scope": "shared",
12
  "n_samples": 128,
 
13
  "symmetric": false,
14
  "use_error_compensation": true,
15
  "use_optimal_shrinkage": true,
@@ -5974,20 +5975,18 @@
5974
  "Model": "Qwen/Qwen3.6-27B",
5975
  "ModelSizeGB": 16.513800712,
5976
  "OriginalSizeGB": 55.5630064,
5977
- "PPL": 7.762491226196289,
5978
- "Params": "{'bits': 3, 'group_size': 128, 'n_samples': 128, 'moe_hessian_scope': 'shared', 'use_error_compensation': True, 'use_optimal_shrinkage': True, 'use_weighted_quantization': True, 'symmetric': False, 'low_memory_optimization': False}",
5979
- "QuantTime": 748.9153661727905,
5980
- "arc_challenge": 59.81228668941979,
5981
- "arc_easy": 77.10437710437711,
5982
- "boolq": 81.55963302752293,
5983
- "commonsense_qa": 84.84848484848484,
5984
- "hellaswag": 82.26448914558853,
5985
- "lambada_openai": 75.66466136231321,
5986
- "openbookqa": 45.4,
5987
- "piqa": 81.99129488574538,
5988
- "social_iqa": 73.0749354005168,
5989
- "truthfulqa_mc2": 54.72186573353546,
5990
- "winogrande": 76.71665351223362,
5991
- "zeroshot_avg": 72.10533470088525
5992
  }
5993
- }
 
10
  "low_memory_optimization": false,
11
  "moe_hessian_scope": "shared",
12
  "n_samples": 128,
13
+ "scale_zero_dtype": "float16",
14
  "symmetric": false,
15
  "use_error_compensation": true,
16
  "use_optimal_shrinkage": true,
 
5975
  "Model": "Qwen/Qwen3.6-27B",
5976
  "ModelSizeGB": 16.513800712,
5977
  "OriginalSizeGB": 55.5630064,
5978
+ "PPL": 7.755378246307373,
5979
+ "Params": "{'bits': 3, 'group_size': 128, 'scale_zero_dtype': 'float16', 'n_samples': 128, 'moe_hessian_scope': 'shared', 'use_error_compensation': True, 'use_optimal_shrinkage': True, 'use_weighted_quantization': True, 'symmetric': False, 'low_memory_optimization': False}",
5980
+ "QuantTime": 787.927401304245,
5981
+ "arc_challenge": 58.78839590443686,
5982
+ "arc_easy": 75.54713804713805,
5983
+ "commonsense_qa": 86.15888615888616,
5984
+ "hellaswag": 82.59310894244175,
5985
+ "lambada_openai": 74.77197748884146,
5986
+ "openbookqa": 44.800000000000004,
5987
+ "piqa": 82.4265505984766,
5988
+ "truthfulqa_mc2": 54.909605261997015,
5989
+ "winogrande": 76.87450670876085,
5990
+ "zeroshot_avg": 70.76335212344208
 
 
5991
  }
5992
+ }
model-00002-of-00004.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:52d1bea86c0ea39d0a6eeb26d86e492089bbaffa3e37dcd0ade1db68c12ee8cb
3
  size 4970468528
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:50ccca36e3cedf1d5cbc3f233f1def041f7502c78ae17b133986212d3bdab0e6
3
  size 4970468528
model-00003-of-00004.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:ce163441bff73002f3fd4df78ebe90cd85e1032376a1920718100b9103ef2ce9
3
  size 4997537440
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:14c351d0bdd2510dc0277ea8ec8294cfb80695560daed88e44ad59f97f69e664
3
  size 4997537440
model-00004-of-00004.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:1c967355c32b68aa0fb2766a790e8362671c2c3a5efa58f2a6531b2facc93526
3
  size 4002997816
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:865ffdcd186d481d4e2a7117bebc88f055c47334acab3ee46fd80904d247a869
3
  size 4002997816