Download TRAINING_PROVENANCE.json from vllm-sr/Decision-1.0-Lex-0.6B: direct link, hf CLI and curl.
- Browser
- Download file 2.31 kB
-
https://huggingface.co/vllm-sr/Decision-1.0-Lex-0.6B/resolve/main/TRAINING_PROVENANCE.json
- Command line
-
hf download hf://vllm-sr/Decision-1.0-Lex-0.6B/TRAINING_PROVENANCE.json
-
curl -L -o TRAINING_PROVENANCE.json https://huggingface.co/vllm-sr/Decision-1.0-Lex-0.6B/resolve/main/TRAINING_PROVENANCE.json
2.31 kB
| { | |
| "active_path_tensors": 486, | |
| "dataset": { | |
| "TEST_used_for_training_or_selection": false, | |
| "cases": 1200, | |
| "decision_exposures": 48000, | |
| "decisions": 6000, | |
| "epochs": 8, | |
| "original_split": "train", | |
| "parquet_bytes": 598824, | |
| "parquet_sha256": "46a58d63edfd86e23229c78afe8b72307bb4ca9fb0e8df180cabb3c67ec9dcd5", | |
| "repo": "LocalLLaMA/typed-decisions", | |
| "revision": "ea9306458d6e9563628369a3d1e72e362fb381d2" | |
| }, | |
| "fresh_native_parity": { | |
| "hard_flip_atol": 0, | |
| "logit_atol": 0.0001, | |
| "max_logit_difference": 0.0, | |
| "max_probability_difference": 0.0, | |
| "official_hard_flips": 0, | |
| "pass": true, | |
| "probability_atol": 2e-05, | |
| "rows": 64 | |
| }, | |
| "frozen_training_plan_sha256": "d15c3cafd15d0ae14cf45e196d987016b304702f6f3930081b2ef30bbeabc638", | |
| "language_scope": "English typed-decisions specialist", | |
| "model": "Decision-1.0-Lex", | |
| "native_files_unchanged": true, | |
| "native_manifest_sha256": "f288d873999832a3f37c6a7c4268c2ab309691e621794dbf7acab891acbbb7e6", | |
| "parameter_tensors": 489, | |
| "parameters": 571909635, | |
| "parent_native_manifest_sha256": "da603662bc57e89ccfb51c972ed9c1f2825f267597353cf1337df9117a3dfabe", | |
| "provenance_scope": "Lex refit; not the Kai tensor-composition ledger or Kai evaluation/performance", | |
| "python_version": "3.12.13", | |
| "recipe": { | |
| "DEV_evaluation_during_final_refit": false, | |
| "bf16_autocast": true, | |
| "clip_norm": 1.0, | |
| "cosine_end_lr": 1e-06, | |
| "encoder_lr": 2.5e-05, | |
| "fixed_export_step": 750, | |
| "head_lr": 0.0001, | |
| "logical_batch_size": 64, | |
| "loss": "cross_entropy(0.5*normalized_original_soft + 0.5*onehot_original_hard)", | |
| "optimizer": "AdamW", | |
| "parameter_precision": "fp32", | |
| "physical_batch_size": 8, | |
| "seed": 20260921, | |
| "steps": 750, | |
| "temperature_fit": false, | |
| "warmup_steps": 0, | |
| "weight_decay": 0.01 | |
| }, | |
| "runtime_versions": { | |
| "numpy": "2.5.3", | |
| "safetensors": "0.8.0", | |
| "tokenizers": "0.22.2", | |
| "torch": "2.12.0+git6bbd260", | |
| "transformers": "4.57.6" | |
| }, | |
| "schema": "decision.lex.training-provenance.v1", | |
| "shared_frozen_tensors": 3, | |
| "training_complete_sha256": "d0cc6eba0cc698cfc6538351ce0f7f73ccd7143be6a9431ad0dd3ae773ba2645", | |
| "training_output_closure_sha256": "d3525627110ea2331d4e32d40303c61f2fd576b8222adfc69394177bc6e28142" | |
| } | |