sandkoan's picture
Upload LoRA adapter from Together AI fine-tuning job ft-021a95e4-5a88
d2b2982 verified
Raw
History Blame Contribute Delete
8.24 kB
{
"best_global_step": null,
"best_metric": null,
"best_model_checkpoint": null,
"epoch": 3.0,
"eval_steps": 0,
"global_step": 45,
"is_hyper_param_search": false,
"is_local_process_zero": true,
"is_world_process_zero": true,
"log_history": [
{
"epoch": 0.06666666666666667,
"grad_norm": 4.894603729248047,
"learning_rate": 1e-05,
"loss": 0.4182,
"step": 1
},
{
"epoch": 0.13333333333333333,
"grad_norm": 5.328320503234863,
"learning_rate": 9.987820251299121e-06,
"loss": 0.464,
"step": 2
},
{
"epoch": 0.2,
"grad_norm": 5.1622443199157715,
"learning_rate": 9.951340343707852e-06,
"loss": 0.4336,
"step": 3
},
{
"epoch": 0.26666666666666666,
"grad_norm": 5.097046375274658,
"learning_rate": 9.890738003669029e-06,
"loss": 0.4309,
"step": 4
},
{
"epoch": 0.3333333333333333,
"grad_norm": 4.911140441894531,
"learning_rate": 9.806308479691595e-06,
"loss": 0.4082,
"step": 5
},
{
"epoch": 0.4,
"grad_norm": 5.099400997161865,
"learning_rate": 9.698463103929542e-06,
"loss": 0.4192,
"step": 6
},
{
"epoch": 0.4666666666666667,
"grad_norm": 4.738089084625244,
"learning_rate": 9.567727288213005e-06,
"loss": 0.3535,
"step": 7
},
{
"epoch": 0.5333333333333333,
"grad_norm": 4.823266983032227,
"learning_rate": 9.414737964294636e-06,
"loss": 0.3764,
"step": 8
},
{
"epoch": 0.6,
"grad_norm": 4.285304546356201,
"learning_rate": 9.24024048078213e-06,
"loss": 0.3005,
"step": 9
},
{
"epoch": 0.6666666666666666,
"grad_norm": 4.088993072509766,
"learning_rate": 9.045084971874738e-06,
"loss": 0.2968,
"step": 10
},
{
"epoch": 0.7333333333333333,
"grad_norm": 4.066262722015381,
"learning_rate": 8.83022221559489e-06,
"loss": 0.3362,
"step": 11
},
{
"epoch": 0.8,
"grad_norm": 3.981386423110962,
"learning_rate": 8.596699001693257e-06,
"loss": 0.2949,
"step": 12
},
{
"epoch": 0.8666666666666667,
"grad_norm": 3.693369150161743,
"learning_rate": 8.345653031794292e-06,
"loss": 0.2606,
"step": 13
},
{
"epoch": 0.9333333333333333,
"grad_norm": 3.731173276901245,
"learning_rate": 8.078307376628292e-06,
"loss": 0.2694,
"step": 14
},
{
"epoch": 1.0,
"grad_norm": 3.457310438156128,
"learning_rate": 7.795964517353734e-06,
"loss": 0.224,
"step": 15
},
{
"epoch": 1.0666666666666667,
"grad_norm": 3.3868319988250732,
"learning_rate": 7.500000000000001e-06,
"loss": 0.2376,
"step": 16
},
{
"epoch": 1.1333333333333333,
"grad_norm": 3.042154550552368,
"learning_rate": 7.191855733945388e-06,
"loss": 0.1872,
"step": 17
},
{
"epoch": 1.2,
"grad_norm": 2.8533551692962646,
"learning_rate": 6.873032967079562e-06,
"loss": 0.1902,
"step": 18
},
{
"epoch": 1.2666666666666666,
"grad_norm": 2.7871158123016357,
"learning_rate": 6.545084971874738e-06,
"loss": 0.1971,
"step": 19
},
{
"epoch": 1.3333333333333333,
"grad_norm": 2.8030219078063965,
"learning_rate": 6.209609477998339e-06,
"loss": 0.206,
"step": 20
},
{
"epoch": 1.4,
"grad_norm": 2.5936150550842285,
"learning_rate": 5.8682408883346535e-06,
"loss": 0.179,
"step": 21
},
{
"epoch": 1.4666666666666668,
"grad_norm": 2.451712131500244,
"learning_rate": 5.522642316338268e-06,
"loss": 0.155,
"step": 22
},
{
"epoch": 1.5333333333333332,
"grad_norm": 2.556074857711792,
"learning_rate": 5.174497483512506e-06,
"loss": 0.1569,
"step": 23
},
{
"epoch": 1.6,
"grad_norm": 2.4171063899993896,
"learning_rate": 4.825502516487497e-06,
"loss": 0.1597,
"step": 24
},
{
"epoch": 1.6666666666666665,
"grad_norm": 2.277461528778076,
"learning_rate": 4.477357683661734e-06,
"loss": 0.1488,
"step": 25
},
{
"epoch": 1.7333333333333334,
"grad_norm": 2.0196449756622314,
"learning_rate": 4.131759111665349e-06,
"loss": 0.1173,
"step": 26
},
{
"epoch": 1.8,
"grad_norm": 2.131396532058716,
"learning_rate": 3.790390522001662e-06,
"loss": 0.1598,
"step": 27
},
{
"epoch": 1.8666666666666667,
"grad_norm": 1.9776694774627686,
"learning_rate": 3.4549150281252635e-06,
"loss": 0.1376,
"step": 28
},
{
"epoch": 1.9333333333333333,
"grad_norm": 1.8695093393325806,
"learning_rate": 3.12696703292044e-06,
"loss": 0.129,
"step": 29
},
{
"epoch": 2.0,
"grad_norm": 1.8652958869934082,
"learning_rate": 2.8081442660546126e-06,
"loss": 0.1298,
"step": 30
},
{
"epoch": 2.066666666666667,
"grad_norm": 1.8152166604995728,
"learning_rate": 2.5000000000000015e-06,
"loss": 0.1277,
"step": 31
},
{
"epoch": 2.1333333333333333,
"grad_norm": 1.745728611946106,
"learning_rate": 2.204035482646267e-06,
"loss": 0.1308,
"step": 32
},
{
"epoch": 2.2,
"grad_norm": 1.7527920007705688,
"learning_rate": 1.9216926233717087e-06,
"loss": 0.1211,
"step": 33
},
{
"epoch": 2.2666666666666666,
"grad_norm": 1.695560097694397,
"learning_rate": 1.6543469682057105e-06,
"loss": 0.1067,
"step": 34
},
{
"epoch": 2.3333333333333335,
"grad_norm": 1.607631802558899,
"learning_rate": 1.4033009983067454e-06,
"loss": 0.0992,
"step": 35
},
{
"epoch": 2.4,
"grad_norm": 1.5756640434265137,
"learning_rate": 1.1697777844051105e-06,
"loss": 0.109,
"step": 36
},
{
"epoch": 2.466666666666667,
"grad_norm": 1.8147677183151245,
"learning_rate": 9.549150281252633e-07,
"loss": 0.1228,
"step": 37
},
{
"epoch": 2.533333333333333,
"grad_norm": 1.5649832487106323,
"learning_rate": 7.597595192178702e-07,
"loss": 0.1134,
"step": 38
},
{
"epoch": 2.6,
"grad_norm": 1.5064197778701782,
"learning_rate": 5.852620357053651e-07,
"loss": 0.1106,
"step": 39
},
{
"epoch": 2.6666666666666665,
"grad_norm": 1.5975162982940674,
"learning_rate": 4.322727117869951e-07,
"loss": 0.0989,
"step": 40
},
{
"epoch": 2.7333333333333334,
"grad_norm": 1.6186599731445312,
"learning_rate": 3.015368960704584e-07,
"loss": 0.1303,
"step": 41
},
{
"epoch": 2.8,
"grad_norm": 1.5868901014328003,
"learning_rate": 1.9369152030840553e-07,
"loss": 0.1119,
"step": 42
},
{
"epoch": 2.8666666666666667,
"grad_norm": 1.5993646383285522,
"learning_rate": 1.0926199633097156e-07,
"loss": 0.1117,
"step": 43
},
{
"epoch": 2.9333333333333336,
"grad_norm": 1.6658110618591309,
"learning_rate": 4.865965629214819e-08,
"loss": 0.1286,
"step": 44
},
{
"epoch": 3.0,
"grad_norm": 1.3958430290222168,
"learning_rate": 1.2179748700879013e-08,
"loss": 0.0948,
"step": 45
}
],
"logging_steps": 1.0,
"max_steps": 45,
"num_input_tokens_seen": 0,
"num_train_epochs": 3,
"save_steps": 0,
"stateful_callbacks": {
"TrainerControl": {
"args": {
"should_epoch_stop": false,
"should_evaluate": false,
"should_log": false,
"should_save": true,
"should_training_stop": true
},
"attributes": {}
}
},
"total_flos": 8.578055705670451e+16,
"train_batch_size": 1,
"trial_name": null,
"trial_params": null
}