Download quantization_report.json from lokinfey/Qwen3_5_4B_Hmm_ONNX: direct link, hf CLI and curl.
- Browser
- Download file 158 kB
-
https://huggingface.co/lokinfey/Qwen3_5_4B_Hmm_ONNX/resolve/main/quantization_report.json
- Command line
-
hf download hf://lokinfey/Qwen3_5_4B_Hmm_ONNX/quantization_report.json
-
curl -L -o quantization_report.json https://huggingface.co/lokinfey/Qwen3_5_4B_Hmm_ONNX/resolve/main/quantization_report.json
158 kB
| { | |
| "compute_capability": "Packed storage may be consumed by native MatMulNBits/GatherBlockQuantized/BlockQuantizedMatMul implementations. MatMulNBits may instead be inlined as nibble BitShift/BitwiseAnd, DequantizeLinear, and float MatMul; this does not change packed storage and makes no promise about the 'cpu' runtime kernel.", | |
| "compute_mode": "runtime-dependent native custom op or inline standard-ONNX fallback", | |
| "converted_from": "Q4_K_M-like mixed GGUF", | |
| "dispositions": [ | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "qtypes": [ | |
| "Q4_K", | |
| "Q6_K" | |
| ], | |
| "source_bytes": 2693990400, | |
| "tensor_count": 249 | |
| }, | |
| { | |
| "disposition": "source float", | |
| "qtypes": [ | |
| "F32" | |
| ], | |
| "source_bytes": 3846144, | |
| "tensor_count": 177 | |
| } | |
| ], | |
| "explicit_float_tensors": [ | |
| { | |
| "disposition": "source float", | |
| "name": "blk.0.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.0.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.0.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.0.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.0.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.0.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.1.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.1.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.1.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.1.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.1.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.1.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.10.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.10.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.10.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.10.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.10.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.10.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.11.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.11.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.11.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.11.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.12.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.12.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.12.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.12.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.12.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.12.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.13.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.13.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.13.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.13.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.13.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.13.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.14.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.14.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.14.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.14.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.14.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.14.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.15.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.15.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.15.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.15.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.16.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.16.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.16.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.16.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.16.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.16.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.17.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.17.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.17.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.17.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.17.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.17.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.18.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.18.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.18.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.18.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.18.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.18.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.19.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.19.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.19.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.19.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.2.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.2.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.2.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.2.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.2.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.2.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.20.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.20.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.20.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.20.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.20.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.20.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.21.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.21.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.21.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.21.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.21.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.21.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.22.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.22.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.22.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.22.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.22.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.22.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.23.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.23.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.23.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.23.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.24.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.24.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.24.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.24.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.24.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.24.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.25.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.25.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.25.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.25.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.25.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.25.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.26.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.26.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.26.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.26.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.26.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.26.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.27.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.27.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.27.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.27.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.28.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.28.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.28.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.28.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.28.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.28.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.29.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.29.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.29.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.29.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.29.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.29.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.3.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.3.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.3.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.3.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.30.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.30.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.30.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.30.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.30.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.30.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.31.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.31.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.31.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.31.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.4.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.4.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.4.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.4.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.4.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.4.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.5.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.5.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.5.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.5.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.5.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.5.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.6.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.6.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.6.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.6.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.6.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.6.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.7.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.7.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.7.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.7.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.8.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.8.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.8.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.8.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.8.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.8.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.9.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.9.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.9.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.9.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.9.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.9.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "output_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| } | |
| ], | |
| "schema_version": 1, | |
| "source_fidelity": false, | |
| "source_qtype_census": [ | |
| { | |
| "qtype": "F32", | |
| "source_bytes": 3846144, | |
| "tensor_count": 177 | |
| }, | |
| { | |
| "qtype": "Q4_K", | |
| "source_bytes": 1647820800, | |
| "tensor_count": 216 | |
| }, | |
| { | |
| "qtype": "Q6_K", | |
| "source_bytes": 1046169600, | |
| "tensor_count": 33 | |
| } | |
| ], | |
| "storage_quantized": true, | |
| "target_storage_format": "INT4 affine block-32", | |
| "tensor_records": [ | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.0.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.0.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.0.attn_qkv.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 17203200, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.0.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.0.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.0.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.0.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.0.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.0.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.0.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.0.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.0.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.0.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.0.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.1.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.1.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.1.attn_qkv.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 17203200, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.1.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.1.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.1.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.1.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.1.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.1.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.1.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.1.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.1.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.1.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.1.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.10.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.10.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.10.attn_qkv.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.10.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.10.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.10.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.10.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.10.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.10.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.10.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.10.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.10.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.10.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.10.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.11.attn_k.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 1474560, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.11.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.11.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.11.attn_output.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.11.attn_q.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.11.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.11.attn_v.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 1474560, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.11.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.11.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.11.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.11.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.12.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.12.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.12.attn_qkv.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 17203200, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.12.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.12.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.12.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.12.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.12.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.12.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.12.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.12.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.12.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.12.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.12.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.13.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.13.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.13.attn_qkv.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.13.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.13.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.13.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.13.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.13.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.13.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.13.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.13.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.13.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.13.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.13.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.14.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.14.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.14.attn_qkv.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.14.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.14.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.14.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.14.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.14.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.14.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.14.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.14.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.14.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.14.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.14.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.15.attn_k.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 1474560, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.15.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.15.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.15.attn_output.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.15.attn_q.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.15.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.15.attn_v.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 2150400, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.15.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.15.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.15.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.15.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.16.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.16.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.16.attn_qkv.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.16.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.16.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.16.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.16.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.16.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.16.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.16.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.16.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.16.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.16.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.16.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.17.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.17.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.17.attn_qkv.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.17.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.17.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.17.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.17.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.17.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.17.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.17.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.17.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.17.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.17.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.17.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.18.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.18.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.18.attn_qkv.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 17203200, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.18.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.18.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.18.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.18.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.18.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.18.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.18.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.18.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.18.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.18.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.18.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.19.attn_k.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 1474560, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.19.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.19.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.19.attn_output.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.19.attn_q.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.19.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.19.attn_v.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 1474560, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.19.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.19.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.19.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.19.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.2.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.2.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.2.attn_qkv.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 17203200, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.2.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.2.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.2.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.2.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.2.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.2.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.2.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.2.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.2.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.2.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.2.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.20.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.20.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.20.attn_qkv.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.20.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.20.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.20.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.20.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.20.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.20.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.20.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.20.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.20.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.20.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.20.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.21.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.21.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.21.attn_qkv.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 17203200, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.21.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.21.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.21.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.21.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.21.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.21.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.21.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.21.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.21.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.21.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.21.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.22.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.22.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.22.attn_qkv.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.22.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.22.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.22.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.22.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.22.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.22.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.22.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.22.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.22.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.22.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.22.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.23.attn_k.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 1474560, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.23.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.23.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.23.attn_output.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.23.attn_q.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.23.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.23.attn_v.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 1474560, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.23.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.23.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.23.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.23.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.24.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.24.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.24.attn_qkv.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 17203200, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.24.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.24.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.24.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.24.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.24.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.24.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.24.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.24.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.24.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.24.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.24.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.25.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.25.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.25.attn_qkv.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.25.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.25.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.25.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.25.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.25.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.25.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.25.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.25.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.25.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.25.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.25.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.26.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.26.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.26.attn_qkv.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.26.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.26.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.26.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.26.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.26.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.26.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.26.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.26.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.26.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.26.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.26.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.27.attn_k.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 1474560, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.27.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.27.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.27.attn_output.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.27.attn_q.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.27.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.27.attn_v.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 2150400, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.27.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.27.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.27.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.27.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.28.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.28.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.28.attn_qkv.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 17203200, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.28.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.28.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.28.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.28.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.28.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.28.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.28.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.28.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.28.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.28.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.28.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.29.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.29.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.29.attn_qkv.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 17203200, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.29.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.29.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.29.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.29.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.29.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.29.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.29.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.29.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.29.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.29.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.29.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.3.attn_k.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 1474560, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.3.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.3.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.3.attn_output.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.3.attn_q.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.3.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.3.attn_v.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 2150400, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.3.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.3.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.3.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.3.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.30.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.30.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.30.attn_qkv.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 17203200, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.30.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.30.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.30.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.30.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.30.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.30.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.30.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.30.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.30.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.30.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.30.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.31.attn_k.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 1474560, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.31.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.31.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.31.attn_output.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.31.attn_q.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.31.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.31.attn_v.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 2150400, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.31.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.31.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.31.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.31.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.4.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.4.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.4.attn_qkv.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.4.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.4.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.4.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.4.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.4.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.4.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.4.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.4.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.4.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.4.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.4.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.5.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.5.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.5.attn_qkv.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.5.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.5.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.5.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.5.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.5.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.5.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.5.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.5.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.5.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.5.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.5.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.6.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.6.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.6.attn_qkv.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 17203200, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.6.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.6.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.6.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.6.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.6.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.6.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.6.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.6.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.6.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.6.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.6.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.7.attn_k.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 1474560, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.7.attn_k_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.7.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.7.attn_output.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.7.attn_q.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.7.attn_q_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 1024, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.7.attn_v.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 1474560, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.7.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.7.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.7.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.7.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.8.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.8.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.8.attn_qkv.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 11796480, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.8.ffn_down.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.8.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.8.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.8.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.8.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.8.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.8.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.8.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.8.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.8.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.8.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.9.attn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.9.attn_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.9.attn_qkv.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 17203200, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.9.ffn_down.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 19353600, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.9.ffn_gate.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.9.ffn_up.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 13271040, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.9.post_attention_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.9.ssm_a", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.9.ssm_alpha.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.9.ssm_beta.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 46080, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.9.ssm_conv1d.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 131072, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.9.ssm_dt.bias", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 128, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "blk.9.ssm_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 512, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "blk.9.ssm_out.weight", | |
| "qtype": "Q4_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 5898240, | |
| "target_storage": "INT4 affine block-32" | |
| }, | |
| { | |
| "disposition": "source float", | |
| "name": "output_norm.weight", | |
| "qtype": "F32", | |
| "reason": "The GGUF tensor is already stored as float.", | |
| "source_bytes": 10240, | |
| "target_storage": "float" | |
| }, | |
| { | |
| "disposition": "lossy dequantize+requantize", | |
| "name": "token_embd.weight", | |
| "qtype": "Q6_K", | |
| "reason": "The source is converted through the declared affine/runtime target.", | |
| "source_bytes": 521472000, | |
| "target_storage": "INT4 affine block-32" | |
| } | |
| ] | |
| } | |