Qwen3.6-35B-A3B-AntiLoop-NVFP4 / mlx_conversion_manifest.json
N8Programs's picture
Add MLX-VLM runtime, configuration, and model card
c179c12 verified
Raw
History Blame
1.01 kB
{
"converter": "convert_qwen36_modelopt_hybrid_to_mlx.py",
"created_at": "2026-07-10T03:37:44.758817+00:00",
"input_scales_dropped": 30971,
"multimodal_upgrade_at": "2026-07-10T04:02:39.296349+00:00",
"norm_weights_shifted": true,
"num_experts": 256,
"num_layers": 40,
"output": "mlx-community/Qwen3.6-35B-A3B-AntiLoop-NVFP4",
"output_shards": 42,
"output_tensor_count": 1808,
"output_total_size_bytes": 21758152396,
"quantized_module_counts": {
"scaled_mxfp8": 130,
"scaled_nvfp4": 121,
"scaled_nvfp4_switch": 120
},
"runtime_file": "modeling_mlx_qwen36_modelopt_hybrid.py",
"source": "N8Programs/Qwen3.6-35B-A3B-AntiLoop-NVFP4@1fc377564024dce4e8e7f2bdc04d34cd869f928f",
"source_tensor_count": 124468,
"vision_shard": "model-vision.safetensors",
"vision_tensor_count": 333,
"vision_tensor_data_bytes": 893142496,
"vision_weight_bytes_requantized": false,
"vlm_runtime_file": "modeling_mlx_vlm_qwen36_modelopt_hybrid.py",
"weight_bytes_requantized": false
}