Razavipour commited on
Commit
829b529
·
verified ·
1 Parent(s): aca4c11

End of training

Browse files
Files changed (4) hide show
  1. README.md +3 -1
  2. all_results.json +6 -6
  3. train_results.json +6 -6
  4. trainer_state.json +30 -16
README.md CHANGED
@@ -3,6 +3,8 @@ library_name: peft
3
  license: cc-by-nc-4.0
4
  base_model: facebook/musicgen-melody
5
  tags:
 
 
6
  - generated_from_trainer
7
  model-index:
8
  - name: musicgen-persian-finetuned_setar_with_meta
@@ -14,7 +16,7 @@ should probably proofread and complete it, then remove this comment. -->
14
 
15
  # musicgen-persian-finetuned_setar_with_meta
16
 
17
- This model is a fine-tuned version of [facebook/musicgen-melody](https://huggingface.co/facebook/musicgen-melody) on an unknown dataset.
18
 
19
  ## Model description
20
 
 
3
  license: cc-by-nc-4.0
4
  base_model: facebook/musicgen-melody
5
  tags:
6
+ - text-to-audio
7
+ - Razavipour/persian-solo-setar_test
8
  - generated_from_trainer
9
  model-index:
10
  - name: musicgen-persian-finetuned_setar_with_meta
 
16
 
17
  # musicgen-persian-finetuned_setar_with_meta
18
 
19
+ This model is a fine-tuned version of [facebook/musicgen-melody](https://huggingface.co/facebook/musicgen-melody) on the RAZAVIPOUR/PERSIAN-SOLO-SETAR_TEST - DEFAULT dataset.
20
 
21
  ## Model description
22
 
all_results.json CHANGED
@@ -1,9 +1,9 @@
1
  {
2
- "epoch": 2.0,
3
- "total_flos": 3577272745008.0,
4
- "train_loss": 36.21625518798828,
5
- "train_runtime": 24.3225,
6
  "train_samples": 8,
7
- "train_samples_per_second": 0.658,
8
- "train_steps_per_second": 0.082
9
  }
 
1
  {
2
+ "epoch": 4.0,
3
+ "total_flos": 7191424590480.0,
4
+ "train_loss": 35.26272487640381,
5
+ "train_runtime": 42.358,
6
  "train_samples": 8,
7
+ "train_samples_per_second": 0.755,
8
+ "train_steps_per_second": 0.094
9
  }
train_results.json CHANGED
@@ -1,9 +1,9 @@
1
  {
2
- "epoch": 2.0,
3
- "total_flos": 3577272745008.0,
4
- "train_loss": 36.21625518798828,
5
- "train_runtime": 24.3225,
6
  "train_samples": 8,
7
- "train_samples_per_second": 0.658,
8
- "train_steps_per_second": 0.082
9
  }
 
1
  {
2
+ "epoch": 4.0,
3
+ "total_flos": 7191424590480.0,
4
+ "train_loss": 35.26272487640381,
5
+ "train_runtime": 42.358,
6
  "train_samples": 8,
7
+ "train_samples_per_second": 0.755,
8
+ "train_steps_per_second": 0.094
9
  }
trainer_state.json CHANGED
@@ -2,41 +2,55 @@
2
  "best_global_step": null,
3
  "best_metric": null,
4
  "best_model_checkpoint": null,
5
- "epoch": 2.0,
6
  "eval_steps": 500,
7
- "global_step": 2,
8
  "is_hyper_param_search": false,
9
  "is_local_process_zero": true,
10
  "is_world_process_zero": true,
11
  "log_history": [
12
  {
13
  "epoch": 1.0,
14
- "grad_norm": 7.861708641052246,
15
  "learning_rate": 0.0002,
16
  "loss": 36.484,
17
  "step": 1
18
  },
19
  {
20
  "epoch": 2.0,
21
- "grad_norm": 8.914444923400879,
22
- "learning_rate": 0.0001,
23
- "loss": 35.9485,
24
  "step": 2
25
  },
26
  {
27
- "epoch": 2.0,
28
- "step": 2,
29
- "total_flos": 3577272745008.0,
30
- "train_loss": 36.21625518798828,
31
- "train_runtime": 24.3225,
32
- "train_samples_per_second": 0.658,
33
- "train_steps_per_second": 0.082
 
 
 
 
 
 
 
 
 
 
 
 
 
 
34
  }
35
  ],
36
  "logging_steps": 1.0,
37
- "max_steps": 2,
38
  "num_input_tokens_seen": 0,
39
- "num_train_epochs": 2,
40
  "save_steps": 500,
41
  "stateful_callbacks": {
42
  "TrainerControl": {
@@ -50,7 +64,7 @@
50
  "attributes": {}
51
  }
52
  },
53
- "total_flos": 3577272745008.0,
54
  "train_batch_size": 2,
55
  "trial_name": null,
56
  "trial_params": null
 
2
  "best_global_step": null,
3
  "best_metric": null,
4
  "best_model_checkpoint": null,
5
+ "epoch": 4.0,
6
  "eval_steps": 500,
7
+ "global_step": 4,
8
  "is_hyper_param_search": false,
9
  "is_local_process_zero": true,
10
  "is_world_process_zero": true,
11
  "log_history": [
12
  {
13
  "epoch": 1.0,
14
+ "grad_norm": 7.8616437911987305,
15
  "learning_rate": 0.0002,
16
  "loss": 36.484,
17
  "step": 1
18
  },
19
  {
20
  "epoch": 2.0,
21
+ "grad_norm": 8.914605140686035,
22
+ "learning_rate": 0.00015000000000000001,
23
+ "loss": 35.9482,
24
  "step": 2
25
  },
26
  {
27
+ "epoch": 3.0,
28
+ "grad_norm": 10.631148338317871,
29
+ "learning_rate": 0.0001,
30
+ "loss": 34.8575,
31
+ "step": 3
32
+ },
33
+ {
34
+ "epoch": 4.0,
35
+ "grad_norm": 11.890551567077637,
36
+ "learning_rate": 5e-05,
37
+ "loss": 33.7612,
38
+ "step": 4
39
+ },
40
+ {
41
+ "epoch": 4.0,
42
+ "step": 4,
43
+ "total_flos": 7191424590480.0,
44
+ "train_loss": 35.26272487640381,
45
+ "train_runtime": 42.358,
46
+ "train_samples_per_second": 0.755,
47
+ "train_steps_per_second": 0.094
48
  }
49
  ],
50
  "logging_steps": 1.0,
51
+ "max_steps": 4,
52
  "num_input_tokens_seen": 0,
53
+ "num_train_epochs": 4,
54
  "save_steps": 500,
55
  "stateful_callbacks": {
56
  "TrainerControl": {
 
64
  "attributes": {}
65
  }
66
  },
67
+ "total_flos": 7191424590480.0,
68
  "train_batch_size": 2,
69
  "trial_name": null,
70
  "trial_params": null