Razavipour commited on
Commit
f6b777e
·
verified ·
1 Parent(s): 0afc273

End of training

Browse files
Files changed (4) hide show
  1. README.md +3 -1
  2. all_results.json +6 -6
  3. train_results.json +6 -6
  4. trainer_state.json +18 -32
README.md CHANGED
@@ -3,6 +3,8 @@ library_name: peft
3
  license: cc-by-nc-4.0
4
  base_model: facebook/musicgen-melody
5
  tags:
 
 
6
  - generated_from_trainer
7
  model-index:
8
  - name: musicgen-persian-finetuned_setar_with_meta
@@ -14,7 +16,7 @@ should probably proofread and complete it, then remove this comment. -->
14
 
15
  # musicgen-persian-finetuned_setar_with_meta
16
 
17
- This model is a fine-tuned version of [facebook/musicgen-melody](https://huggingface.co/facebook/musicgen-melody) on an unknown dataset.
18
 
19
  ## Model description
20
 
 
3
  license: cc-by-nc-4.0
4
  base_model: facebook/musicgen-melody
5
  tags:
6
+ - text-to-audio
7
+ - Razavipour/persian-solo-setar_test
8
  - generated_from_trainer
9
  model-index:
10
  - name: musicgen-persian-finetuned_setar_with_meta
 
16
 
17
  # musicgen-persian-finetuned_setar_with_meta
18
 
19
+ This model is a fine-tuned version of [facebook/musicgen-melody](https://huggingface.co/facebook/musicgen-melody) on the RAZAVIPOUR/PERSIAN-SOLO-SETAR_TEST - DEFAULT dataset.
20
 
21
  ## Model description
22
 
all_results.json CHANGED
@@ -1,9 +1,9 @@
1
  {
2
- "epoch": 4.0,
3
- "total_flos": 3245360840832.0,
4
- "train_loss": 17.523850917816162,
5
- "train_runtime": 50.1803,
6
  "train_samples": 8,
7
- "train_samples_per_second": 0.638,
8
- "train_steps_per_second": 0.08
9
  }
 
1
  {
2
+ "epoch": 2.0,
3
+ "total_flos": 3577272745008.0,
4
+ "train_loss": 36.21625518798828,
5
+ "train_runtime": 24.3225,
6
  "train_samples": 8,
7
+ "train_samples_per_second": 0.658,
8
+ "train_steps_per_second": 0.082
9
  }
train_results.json CHANGED
@@ -1,9 +1,9 @@
1
  {
2
- "epoch": 4.0,
3
- "total_flos": 3245360840832.0,
4
- "train_loss": 17.523850917816162,
5
- "train_runtime": 50.1803,
6
  "train_samples": 8,
7
- "train_samples_per_second": 0.638,
8
- "train_steps_per_second": 0.08
9
  }
 
1
  {
2
+ "epoch": 2.0,
3
+ "total_flos": 3577272745008.0,
4
+ "train_loss": 36.21625518798828,
5
+ "train_runtime": 24.3225,
6
  "train_samples": 8,
7
+ "train_samples_per_second": 0.658,
8
+ "train_steps_per_second": 0.082
9
  }
trainer_state.json CHANGED
@@ -2,55 +2,41 @@
2
  "best_global_step": null,
3
  "best_metric": null,
4
  "best_model_checkpoint": null,
5
- "epoch": 4.0,
6
  "eval_steps": 500,
7
- "global_step": 4,
8
  "is_hyper_param_search": false,
9
  "is_local_process_zero": true,
10
  "is_world_process_zero": true,
11
  "log_history": [
12
  {
13
  "epoch": 1.0,
14
- "grad_norm": 4.468116760253906,
15
  "learning_rate": 0.0002,
16
- "loss": 18.2546,
17
  "step": 1
18
  },
19
  {
20
  "epoch": 2.0,
21
- "grad_norm": 4.584405899047852,
22
- "learning_rate": 0.00015000000000000001,
23
- "loss": 17.7667,
24
- "step": 2
25
- },
26
- {
27
- "epoch": 3.0,
28
- "grad_norm": 5.513431549072266,
29
  "learning_rate": 0.0001,
30
- "loss": 17.3348,
31
- "step": 3
32
- },
33
- {
34
- "epoch": 4.0,
35
- "grad_norm": 6.110393524169922,
36
- "learning_rate": 5e-05,
37
- "loss": 16.7393,
38
- "step": 4
39
  },
40
  {
41
- "epoch": 4.0,
42
- "step": 4,
43
- "total_flos": 3245360840832.0,
44
- "train_loss": 17.523850917816162,
45
- "train_runtime": 50.1803,
46
- "train_samples_per_second": 0.638,
47
- "train_steps_per_second": 0.08
48
  }
49
  ],
50
  "logging_steps": 1.0,
51
- "max_steps": 4,
52
  "num_input_tokens_seen": 0,
53
- "num_train_epochs": 4,
54
  "save_steps": 500,
55
  "stateful_callbacks": {
56
  "TrainerControl": {
@@ -64,8 +50,8 @@
64
  "attributes": {}
65
  }
66
  },
67
- "total_flos": 3245360840832.0,
68
- "train_batch_size": 4,
69
  "trial_name": null,
70
  "trial_params": null
71
  }
 
2
  "best_global_step": null,
3
  "best_metric": null,
4
  "best_model_checkpoint": null,
5
+ "epoch": 2.0,
6
  "eval_steps": 500,
7
+ "global_step": 2,
8
  "is_hyper_param_search": false,
9
  "is_local_process_zero": true,
10
  "is_world_process_zero": true,
11
  "log_history": [
12
  {
13
  "epoch": 1.0,
14
+ "grad_norm": 7.861708641052246,
15
  "learning_rate": 0.0002,
16
+ "loss": 36.484,
17
  "step": 1
18
  },
19
  {
20
  "epoch": 2.0,
21
+ "grad_norm": 8.914444923400879,
 
 
 
 
 
 
 
22
  "learning_rate": 0.0001,
23
+ "loss": 35.9485,
24
+ "step": 2
 
 
 
 
 
 
 
25
  },
26
  {
27
+ "epoch": 2.0,
28
+ "step": 2,
29
+ "total_flos": 3577272745008.0,
30
+ "train_loss": 36.21625518798828,
31
+ "train_runtime": 24.3225,
32
+ "train_samples_per_second": 0.658,
33
+ "train_steps_per_second": 0.082
34
  }
35
  ],
36
  "logging_steps": 1.0,
37
+ "max_steps": 2,
38
  "num_input_tokens_seen": 0,
39
+ "num_train_epochs": 2,
40
  "save_steps": 500,
41
  "stateful_callbacks": {
42
  "TrainerControl": {
 
50
  "attributes": {}
51
  }
52
  },
53
+ "total_flos": 3577272745008.0,
54
+ "train_batch_size": 2,
55
  "trial_name": null,
56
  "trial_params": null
57
  }