{ "label": "SmolLM3-3B APO LoRA, supervised (published re-run)", "family": "Readout fine-tuning \u00b7 SmolLM3 (after SFT and after APO)", "size": "3B", "stage": "LoRA sup", "acc": 0.7052638906784982, "ece": 0.057626918841681425, "brier": 0.3628097830416841, "heldout_acc": 0.7406916545366327, "heldout_ece": 0.07353221732093843, "trained_acc": 0.6919784792316978, "tvd": 0.33711000389760293, "n_configs": 22, "T_choice": 1.2192, "T_score": 1.4924, "T_noul": 1.8579, "sure_loss": 0.2919306388122165, "frac_incoherent": 0.9641855705270341, "proj_gain_acc": 0.008257433623287236, "proj_gain_ece": 0.009138484289030219, "chaos_rho": 0.18360444399307768 }