{ "format_version": 1, "project": "dohnuts", "distribution": "dohnuts", "base_model": "Qwen/Qwen3.5-0.8B", "base_revision": "2fc06364715b967f1860aea9cf38778875588b17", "adapter": "qwen3.5", "lora_rank": 8, "image_pixels": 262144, "max_length": 2048, "backend": { "linear_patch": true, "triton_convolution": true, "experimental_rocm_sdpa": true, "fused_norm_and_swiglu": true, "shared_prefix": true, "frozen_vision_cache_MiB": 128 }, "inference": "merged LoRA, BF16, fused operations, shared-prefix parallel candidate scoring", "source_hashes": { "train": "22c5ba82b2b9c69a0a9ce9378a83bb50004478031a3d1e2167a973cda881ed94", "dev": "04b3cf11acece13f4bb44842c7a55f11c9c30d6ed99819246f5676d6da6497a9", "calibration": "c96f04a195ef39c996b4c42d49f467e7a2d8bb0ef248a9d865ae5af6ece84482", "test": "30504c618ab7a72a91ff14520d2e79552f47ce16519f1a7d96c64858d7a5ad79" }, "selected_step": 3600, "selection": "maximum development macro accuracy within the run; no test selection", "seed": 42, "dev_macro_accuracy": 0.7783954326923077, "temperatures": { "choice": 1.8172028064727783, "noul": 3.5866143703460693, "score": 1.2567481994628906 }, "weights_sha256": "196be33a0282537bcd821e2115643b352d0ad2a0bbaf7b242a1b7fe5bd96cfdf", "version": "0.1.0", "model_id": "Dohnuts-0.1.0-0.8B" }