abarbosa commited on
Commit
5ad996b
·
verified ·
1 Parent(s): 4581f4e

Pushing fine-tuned model to Hugging Face Hub

Browse files
README.md ADDED
@@ -0,0 +1,48 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+
2
+ ---
3
+ language:
4
+ - pt
5
+ - en
6
+ tags:
7
+ - aes
8
+ datasets:
9
+ - kamel-usp/aes_enem_dataset
10
+ base_model: PORTULAN/albertina-1b5-portuguese-ptbr-encoder
11
+ metrics:
12
+ - accuracy
13
+ - qwk
14
+ library_name: transformers
15
+ model-index:
16
+ - name: albertina-1b5-portuguese-ptbr-encoder-encoder_classification-C1-essay_only
17
+ results:
18
+ - task:
19
+ type: text-classification
20
+ name: Automated Essay Score
21
+ dataset:
22
+ name: Automated Essay Score ENEM Dataset
23
+ type: kamel-usp/aes_enem_dataset
24
+ config: JBCS2025
25
+ split: test
26
+ metrics:
27
+ - name: Macro F1
28
+ type: f1
29
+ value: 0.4968802779669244
30
+ - name: QWK
31
+ type: qwk
32
+ value: 0.6826328310864394
33
+ - name: Weighted Macro F1
34
+ type: f1
35
+ value: 0.7116672747290037
36
+ ---
37
+ # Model ID: albertina-1b5-portuguese-ptbr-encoder-encoder_classification-C1-essay_only
38
+ ## Results
39
+ | | test_data |
40
+ |:-----------------|------------:|
41
+ | eval_accuracy | 0.702899 |
42
+ | eval_RMSE | 25.9319 |
43
+ | eval_QWK | 0.682633 |
44
+ | eval_Macro_F1 | 0.49688 |
45
+ | eval_Weighted_F1 | 0.711667 |
46
+ | eval_Micro_F1 | 0.702899 |
47
+ | eval_HDIV | 0.00724638 |
48
+
config.json ADDED
@@ -0,0 +1,54 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "architectures": [
3
+ "DebertaV2ForSequenceClassification"
4
+ ],
5
+ "attention_head_size": 64,
6
+ "attention_probs_dropout_prob": 0.1,
7
+ "conv_act": "gelu",
8
+ "conv_kernel_size": 3,
9
+ "hidden_act": "gelu",
10
+ "hidden_dropout_prob": 0.1,
11
+ "hidden_size": 1536,
12
+ "id2label": {
13
+ "0": 0,
14
+ "1": 40,
15
+ "2": 80,
16
+ "3": 120,
17
+ "4": 160,
18
+ "5": 200
19
+ },
20
+ "initializer_range": 0.02,
21
+ "intermediate_size": 6144,
22
+ "label2id": {
23
+ "0": 0,
24
+ "40": 1,
25
+ "80": 2,
26
+ "120": 3,
27
+ "160": 4,
28
+ "200": 5
29
+ },
30
+ "layer_norm_eps": 1e-07,
31
+ "legacy": true,
32
+ "max_position_embeddings": 512,
33
+ "max_relative_positions": -1,
34
+ "model_type": "deberta-v2",
35
+ "norm_rel_ebd": "layer_norm",
36
+ "num_attention_heads": 24,
37
+ "num_hidden_layers": 48,
38
+ "pad_token_id": 0,
39
+ "pooler_dropout": 0,
40
+ "pooler_hidden_act": "gelu",
41
+ "pooler_hidden_size": 1536,
42
+ "pos_att_type": [
43
+ "p2c",
44
+ "c2p"
45
+ ],
46
+ "position_biased_input": false,
47
+ "position_buckets": 256,
48
+ "relative_attention": true,
49
+ "share_att_key": true,
50
+ "torch_dtype": "bfloat16",
51
+ "transformers_version": "4.53.1",
52
+ "type_vocab_size": 0,
53
+ "vocab_size": 128100
54
+ }
emissions.csv ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ timestamp,project_name,run_id,experiment_id,duration,emissions,emissions_rate,cpu_power,gpu_power,ram_power,cpu_energy,gpu_energy,ram_energy,energy_consumed,country_name,country_iso_code,region,cloud_provider,cloud_region,os,python_version,codecarbon_version,cpu_count,cpu_model,gpu_count,gpu_model,longitude,latitude,ram_total_size,tracking_mode,on_cloud,pue
2
+ 2025-07-10T17:08:59,jbcs2025,7fb4d199-6f0f-4619-b491-03e14cbeafbd,albertina-1b5-portuguese-ptbr-encoder-encoder_classification-C1-essay_only,4549.216668849986,0.10359822916253514,2.2772762148680895e-05,50.27055555555555,239.75622680535508,58.0,0.05155944090119914,0.30703561368385035,0.07201820222251853,0.430613256807568,Romania,ROU,gorj county,,,Linux-5.15.0-143-generic-x86_64-with-glibc2.35,3.12.11,3.0.2,36,Intel(R) Xeon(R) Gold 6248R CPU @ 3.00GHz,1,1 x NVIDIA RTX A6000,23.2904,45.0489,393.6063117980957,machine,N,1.0
evaluation_results.csv ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ eval_loss,eval_model_preparation_time,eval_accuracy,eval_RMSE,eval_QWK,eval_HDIV,eval_Macro_F1,eval_Micro_F1,eval_Weighted_F1,eval_TP_0,eval_TN_0,eval_FP_0,eval_FN_0,eval_TP_1,eval_TN_1,eval_FP_1,eval_FN_1,eval_TP_2,eval_TN_2,eval_FP_2,eval_FN_2,eval_TP_3,eval_TN_3,eval_FP_3,eval_FN_3,eval_TP_4,eval_TN_4,eval_FP_4,eval_FN_4,eval_TP_5,eval_TN_5,eval_FP_5,eval_FN_5,eval_runtime,eval_samples_per_second,eval_steps_per_second,epoch,reference,timestamp,id
2
+ 1.892467737197876,0.0074,0.11363636363636363,99.39209163101397,-0.012651078730374188,0.3787878787878788,0.08164568509396096,0.11363636363636363,0.13984674329501917,0,111,20,1,0,88,44,0,1,112,13,6,0,93,0,39,9,60,9,54,5,79,31,17,8.8626,14.894,3.723,-1,validation_before_training,2025-07-10 16:02:29,albertina-1b5-portuguese-ptbr-encoder-encoder_classification-C1-essay_only
3
+ 1.0570592880249023,0.0074,0.6439393939393939,25.34608929251695,0.6863342898134864,0.0,0.4506612692740875,0.6439393939393939,0.6292028652552237,0,131,0,1,0,132,0,0,4,119,6,3,20,85,8,19,53,42,27,10,8,104,6,14,8.4519,15.618,3.904,9.0,validation_after_training,2025-07-10 16:02:29,albertina-1b5-portuguese-ptbr-encoder-encoder_classification-C1-essay_only
4
+ 0.9303792119026184,0.0074,0.7028985507246377,25.931906372573962,0.6826328310864394,0.007246376811594235,0.49688027796692447,0.7028985507246377,0.7116672747290037,0,137,0,1,0,138,0,0,6,122,6,4,45,65,7,21,40,71,16,11,6,116,12,4,8.7537,15.765,3.998,9.0,test_results,2025-07-10 16:02:29,albertina-1b5-portuguese-ptbr-encoder-encoder_classification-C1-essay_only
model.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:5ccf1f1486296e6656aff978b8d27ed922e9151c1c19c1447996f895e767e052
3
+ size 3133938428
run_experiment.log ADDED
@@ -0,0 +1,305 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ [2025-07-10 15:53:06,842][__main__][INFO] - cache_dir: /tmp/
2
+ dataset:
3
+ name: kamel-usp/aes_enem_dataset
4
+ split: JBCS2025
5
+ training_params:
6
+ seed: 42
7
+ num_train_epochs: 20
8
+ logging_steps: 100
9
+ metric_for_best_model: QWK
10
+ bf16: true
11
+ bootstrap:
12
+ enabled: true
13
+ n_bootstrap: 10000
14
+ bootstrap_seed: 42
15
+ metrics:
16
+ - QWK
17
+ - Macro_F1
18
+ - Weighted_F1
19
+ post_training_results:
20
+ model_path: /workspace/jbcs2025/outputs/2025-03-24/20-42-59
21
+ experiments:
22
+ model:
23
+ name: PORTULAN/albertina-1b5-portuguese-ptbr-encoder
24
+ type: encoder_classification
25
+ num_labels: 6
26
+ output_dir: ./results/
27
+ logging_dir: ./logs/
28
+ best_model_dir: ./results/best_model
29
+ tokenizer:
30
+ name: PORTULAN/albertina-1b5-portuguese-ptbr-encoder
31
+ dataset:
32
+ grade_index: 0
33
+ use_full_context: false
34
+ training_params:
35
+ weight_decay: 0.01
36
+ warmup_ratio: 0.1
37
+ learning_rate: 5.0e-05
38
+ train_batch_size: 4
39
+ eval_batch_size: 4
40
+ gradient_accumulation_steps: 4
41
+ gradient_checkpointing: false
42
+
43
+ [2025-07-10 15:53:10,740][__main__][INFO] - GPU 0: NVIDIA RTX A6000 | TDP ≈ 300 W
44
+ [2025-07-10 15:53:10,740][__main__][INFO] - Starting the Fine Tuning training process.
45
+ [2025-07-10 15:53:15,881][transformers.configuration_utils][INFO] - loading configuration file config.json from cache at /tmp/models--PORTULAN--albertina-1b5-portuguese-ptbr-encoder/snapshots/b22008e5096af9c398b75762d4e28e5008762916/config.json
46
+ [2025-07-10 15:53:15,884][transformers.configuration_utils][INFO] - Model config DebertaV2Config {
47
+ "architectures": [
48
+ "DebertaV2ForMaskedLM"
49
+ ],
50
+ "attention_head_size": 64,
51
+ "attention_probs_dropout_prob": 0.1,
52
+ "conv_act": "gelu",
53
+ "conv_kernel_size": 3,
54
+ "hidden_act": "gelu",
55
+ "hidden_dropout_prob": 0.1,
56
+ "hidden_size": 1536,
57
+ "initializer_range": 0.02,
58
+ "intermediate_size": 6144,
59
+ "layer_norm_eps": 1e-07,
60
+ "legacy": true,
61
+ "max_position_embeddings": 512,
62
+ "max_relative_positions": -1,
63
+ "model_type": "deberta-v2",
64
+ "norm_rel_ebd": "layer_norm",
65
+ "num_attention_heads": 24,
66
+ "num_hidden_layers": 48,
67
+ "pad_token_id": 0,
68
+ "pooler_dropout": 0,
69
+ "pooler_hidden_act": "gelu",
70
+ "pooler_hidden_size": 1536,
71
+ "pos_att_type": [
72
+ "p2c",
73
+ "c2p"
74
+ ],
75
+ "position_biased_input": false,
76
+ "position_buckets": 256,
77
+ "relative_attention": true,
78
+ "share_att_key": true,
79
+ "torch_dtype": "bfloat16",
80
+ "transformers_version": "4.53.1",
81
+ "type_vocab_size": 0,
82
+ "vocab_size": 128100
83
+ }
84
+
85
+ [2025-07-10 15:53:16,306][transformers.tokenization_utils_base][INFO] - loading file spm.model from cache at /tmp/models--PORTULAN--albertina-1b5-portuguese-ptbr-encoder/snapshots/b22008e5096af9c398b75762d4e28e5008762916/spm.model
86
+ [2025-07-10 15:53:16,307][transformers.tokenization_utils_base][INFO] - loading file tokenizer.json from cache at /tmp/models--PORTULAN--albertina-1b5-portuguese-ptbr-encoder/snapshots/b22008e5096af9c398b75762d4e28e5008762916/tokenizer.json
87
+ [2025-07-10 15:53:16,307][transformers.tokenization_utils_base][INFO] - loading file added_tokens.json from cache at /tmp/models--PORTULAN--albertina-1b5-portuguese-ptbr-encoder/snapshots/b22008e5096af9c398b75762d4e28e5008762916/added_tokens.json
88
+ [2025-07-10 15:53:16,307][transformers.tokenization_utils_base][INFO] - loading file special_tokens_map.json from cache at /tmp/models--PORTULAN--albertina-1b5-portuguese-ptbr-encoder/snapshots/b22008e5096af9c398b75762d4e28e5008762916/special_tokens_map.json
89
+ [2025-07-10 15:53:16,307][transformers.tokenization_utils_base][INFO] - loading file tokenizer_config.json from cache at /tmp/models--PORTULAN--albertina-1b5-portuguese-ptbr-encoder/snapshots/b22008e5096af9c398b75762d4e28e5008762916/tokenizer_config.json
90
+ [2025-07-10 15:53:16,307][transformers.tokenization_utils_base][INFO] - loading file chat_template.jinja from cache at None
91
+ [2025-07-10 15:53:16,570][__main__][INFO] - Tokenizer function parameters- Padding:longest; Truncation: True; Use Full Context: False
92
+ [2025-07-10 15:53:16,972][__main__][INFO] -
93
+ Token statistics for 'train' split:
94
+ [2025-07-10 15:53:16,972][__main__][INFO] - Total examples: 500
95
+ [2025-07-10 15:53:16,972][__main__][INFO] - Min tokens: 512
96
+ [2025-07-10 15:53:16,972][__main__][INFO] - Max tokens: 512
97
+ [2025-07-10 15:53:16,972][__main__][INFO] - Avg tokens: 512.00
98
+ [2025-07-10 15:53:16,972][__main__][INFO] - Std tokens: 0.00
99
+ [2025-07-10 15:53:17,067][__main__][INFO] -
100
+ Token statistics for 'validation' split:
101
+ [2025-07-10 15:53:17,067][__main__][INFO] - Total examples: 132
102
+ [2025-07-10 15:53:17,067][__main__][INFO] - Min tokens: 512
103
+ [2025-07-10 15:53:17,067][__main__][INFO] - Max tokens: 512
104
+ [2025-07-10 15:53:17,067][__main__][INFO] - Avg tokens: 512.00
105
+ [2025-07-10 15:53:17,067][__main__][INFO] - Std tokens: 0.00
106
+ [2025-07-10 15:53:17,166][__main__][INFO] -
107
+ Token statistics for 'test' split:
108
+ [2025-07-10 15:53:17,166][__main__][INFO] - Total examples: 138
109
+ [2025-07-10 15:53:17,167][__main__][INFO] - Min tokens: 512
110
+ [2025-07-10 15:53:17,167][__main__][INFO] - Max tokens: 512
111
+ [2025-07-10 15:53:17,167][__main__][INFO] - Avg tokens: 512.00
112
+ [2025-07-10 15:53:17,167][__main__][INFO] - Std tokens: 0.00
113
+ [2025-07-10 15:53:17,167][__main__][INFO] - If token statistics are the same (max, avg, min) keep in mind that this is due to batched tokenization and padding.
114
+ [2025-07-10 15:53:17,167][__main__][INFO] - Model max length: 512. If it is the same as stats, then there is a high chance that sequences are being truncated.
115
+ [2025-07-10 15:53:17,431][transformers.configuration_utils][INFO] - loading configuration file config.json from cache at /tmp/models--PORTULAN--albertina-1b5-portuguese-ptbr-encoder/snapshots/b22008e5096af9c398b75762d4e28e5008762916/config.json
116
+ [2025-07-10 15:53:17,432][transformers.configuration_utils][INFO] - Model config DebertaV2Config {
117
+ "architectures": [
118
+ "DebertaV2ForMaskedLM"
119
+ ],
120
+ "attention_head_size": 64,
121
+ "attention_probs_dropout_prob": 0.1,
122
+ "conv_act": "gelu",
123
+ "conv_kernel_size": 3,
124
+ "hidden_act": "gelu",
125
+ "hidden_dropout_prob": 0.1,
126
+ "hidden_size": 1536,
127
+ "id2label": {
128
+ "0": 0,
129
+ "1": 40,
130
+ "2": 80,
131
+ "3": 120,
132
+ "4": 160,
133
+ "5": 200
134
+ },
135
+ "initializer_range": 0.02,
136
+ "intermediate_size": 6144,
137
+ "label2id": {
138
+ "0": 0,
139
+ "40": 1,
140
+ "80": 2,
141
+ "120": 3,
142
+ "160": 4,
143
+ "200": 5
144
+ },
145
+ "layer_norm_eps": 1e-07,
146
+ "legacy": true,
147
+ "max_position_embeddings": 512,
148
+ "max_relative_positions": -1,
149
+ "model_type": "deberta-v2",
150
+ "norm_rel_ebd": "layer_norm",
151
+ "num_attention_heads": 24,
152
+ "num_hidden_layers": 48,
153
+ "pad_token_id": 0,
154
+ "pooler_dropout": 0,
155
+ "pooler_hidden_act": "gelu",
156
+ "pooler_hidden_size": 1536,
157
+ "pos_att_type": [
158
+ "p2c",
159
+ "c2p"
160
+ ],
161
+ "position_biased_input": false,
162
+ "position_buckets": 256,
163
+ "relative_attention": true,
164
+ "share_att_key": true,
165
+ "torch_dtype": "bfloat16",
166
+ "transformers_version": "4.53.1",
167
+ "type_vocab_size": 0,
168
+ "vocab_size": 128100
169
+ }
170
+
171
+ [2025-07-10 15:53:17,820][transformers.modeling_utils][INFO] - loading weights file pytorch_model.bin from cache at /tmp/models--PORTULAN--albertina-1b5-portuguese-ptbr-encoder/snapshots/b22008e5096af9c398b75762d4e28e5008762916/pytorch_model.bin
172
+ [2025-07-10 15:53:17,821][transformers.modeling_utils][INFO] - Will use torch_dtype=torch.bfloat16 as defined in model's config object
173
+ [2025-07-10 15:53:17,821][transformers.modeling_utils][INFO] - Instantiating DebertaV2ForSequenceClassification model under default dtype torch.bfloat16.
174
+ [2025-07-10 15:53:18,201][transformers.safetensors_conversion][INFO] - Attempting to create safetensors variant
175
+ [2025-07-10 15:53:18,665][transformers.safetensors_conversion][INFO] - Attempting to convert .bin model on the fly to safetensors.
176
+ [2025-07-10 16:02:29,595][transformers.modeling_utils][INFO] - Some weights of the model checkpoint at PORTULAN/albertina-1b5-portuguese-ptbr-encoder were not used when initializing DebertaV2ForSequenceClassification: ['cls.predictions.bias', 'cls.predictions.decoder.bias', 'cls.predictions.decoder.weight', 'cls.predictions.transform.LayerNorm.bias', 'cls.predictions.transform.LayerNorm.weight', 'cls.predictions.transform.dense.bias', 'cls.predictions.transform.dense.weight']
177
+ - This IS expected if you are initializing DebertaV2ForSequenceClassification from the checkpoint of a model trained on another task or with another architecture (e.g. initializing a BertForSequenceClassification model from a BertForPreTraining model).
178
+ - This IS NOT expected if you are initializing DebertaV2ForSequenceClassification from the checkpoint of a model that you expect to be exactly identical (initializing a BertForSequenceClassification model from a BertForSequenceClassification model).
179
+ [2025-07-10 16:02:29,595][transformers.modeling_utils][WARNING] - Some weights of DebertaV2ForSequenceClassification were not initialized from the model checkpoint at PORTULAN/albertina-1b5-portuguese-ptbr-encoder and are newly initialized: ['classifier.bias', 'classifier.weight', 'pooler.dense.bias', 'pooler.dense.weight']
180
+ You should probably TRAIN this model on a down-stream task to be able to use it for predictions and inference.
181
+ [2025-07-10 16:02:29,613][transformers.training_args][INFO] - PyTorch: setting up devices
182
+ [2025-07-10 16:02:29,639][__main__][INFO] - Total steps: 620. Number of warmup steps: 62
183
+ [2025-07-10 16:02:29,649][transformers.trainer][INFO] - You have loaded a model on multiple GPUs. `is_model_parallel` attribute will be force-set to `True` to avoid any unexpected behavior such as device placement mismatching.
184
+ [2025-07-10 16:02:29,676][transformers.trainer][INFO] - Using auto half precision backend
185
+ [2025-07-10 16:02:29,678][transformers.trainer][INFO] - The following columns in the Evaluation set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
186
+ [2025-07-10 16:02:29,689][transformers.trainer][INFO] -
187
+ ***** Running Evaluation *****
188
+ [2025-07-10 16:02:29,689][transformers.trainer][INFO] - Num examples = 132
189
+ [2025-07-10 16:02:29,690][transformers.trainer][INFO] - Batch size = 4
190
+ [2025-07-10 16:02:38,845][transformers.trainer][INFO] - The following columns in the Training set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
191
+ [2025-07-10 16:02:38,876][transformers.trainer][INFO] - ***** Running training *****
192
+ [2025-07-10 16:02:38,876][transformers.trainer][INFO] - Num examples = 500
193
+ [2025-07-10 16:02:38,876][transformers.trainer][INFO] - Num Epochs = 20
194
+ [2025-07-10 16:02:38,876][transformers.trainer][INFO] - Instantaneous batch size per device = 4
195
+ [2025-07-10 16:02:38,876][transformers.trainer][INFO] - Total train batch size (w. parallel, distributed & accumulation) = 16
196
+ [2025-07-10 16:02:38,876][transformers.trainer][INFO] - Gradient Accumulation steps = 4
197
+ [2025-07-10 16:02:38,876][transformers.trainer][INFO] - Total optimization steps = 640
198
+ [2025-07-10 16:02:38,878][transformers.trainer][INFO] - Number of trainable parameters = 1,566,919,686
199
+ [2025-07-10 16:09:32,605][transformers.trainer][INFO] - The following columns in the Evaluation set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
200
+ [2025-07-10 16:09:32,609][transformers.trainer][INFO] -
201
+ ***** Running Evaluation *****
202
+ [2025-07-10 16:09:32,609][transformers.trainer][INFO] - Num examples = 132
203
+ [2025-07-10 16:09:32,609][transformers.trainer][INFO] - Batch size = 4
204
+ [2025-07-10 16:09:41,083][transformers.trainer][INFO] - Saving model checkpoint to /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-32
205
+ [2025-07-10 16:09:41,085][transformers.configuration_utils][INFO] - Configuration saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-32/config.json
206
+ [2025-07-10 16:09:46,899][transformers.modeling_utils][INFO] - Model weights saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-32/model.safetensors
207
+ [2025-07-10 16:16:50,105][transformers.trainer][INFO] - The following columns in the Evaluation set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
208
+ [2025-07-10 16:16:50,109][transformers.trainer][INFO] -
209
+ ***** Running Evaluation *****
210
+ [2025-07-10 16:16:50,110][transformers.trainer][INFO] - Num examples = 132
211
+ [2025-07-10 16:16:50,110][transformers.trainer][INFO] - Batch size = 4
212
+ [2025-07-10 16:16:58,598][transformers.trainer][INFO] - Saving model checkpoint to /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-64
213
+ [2025-07-10 16:16:58,601][transformers.configuration_utils][INFO] - Configuration saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-64/config.json
214
+ [2025-07-10 16:17:05,052][transformers.modeling_utils][INFO] - Model weights saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-64/model.safetensors
215
+ [2025-07-10 16:17:13,233][transformers.trainer][INFO] - Deleting older checkpoint [/workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-32] due to args.save_total_limit
216
+ [2025-07-10 16:24:09,684][transformers.trainer][INFO] - The following columns in the Evaluation set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
217
+ [2025-07-10 16:24:09,688][transformers.trainer][INFO] -
218
+ ***** Running Evaluation *****
219
+ [2025-07-10 16:24:09,688][transformers.trainer][INFO] - Num examples = 132
220
+ [2025-07-10 16:24:09,688][transformers.trainer][INFO] - Batch size = 4
221
+ [2025-07-10 16:24:18,146][transformers.trainer][INFO] - Saving model checkpoint to /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-96
222
+ [2025-07-10 16:24:18,148][transformers.configuration_utils][INFO] - Configuration saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-96/config.json
223
+ [2025-07-10 16:24:23,608][transformers.modeling_utils][INFO] - Model weights saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-96/model.safetensors
224
+ [2025-07-10 16:24:30,971][transformers.trainer][INFO] - Deleting older checkpoint [/workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-64] due to args.save_total_limit
225
+ [2025-07-10 16:31:27,837][transformers.trainer][INFO] - The following columns in the Evaluation set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
226
+ [2025-07-10 16:31:27,841][transformers.trainer][INFO] -
227
+ ***** Running Evaluation *****
228
+ [2025-07-10 16:31:27,842][transformers.trainer][INFO] - Num examples = 132
229
+ [2025-07-10 16:31:27,842][transformers.trainer][INFO] - Batch size = 4
230
+ [2025-07-10 16:31:36,301][transformers.trainer][INFO] - Saving model checkpoint to /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-128
231
+ [2025-07-10 16:31:36,304][transformers.configuration_utils][INFO] - Configuration saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-128/config.json
232
+ [2025-07-10 16:31:43,337][transformers.modeling_utils][INFO] - Model weights saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-128/model.safetensors
233
+ [2025-07-10 16:31:49,743][transformers.trainer][INFO] - Deleting older checkpoint [/workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-96] due to args.save_total_limit
234
+ [2025-07-10 16:38:52,376][transformers.trainer][INFO] - The following columns in the Evaluation set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
235
+ [2025-07-10 16:38:52,380][transformers.trainer][INFO] -
236
+ ***** Running Evaluation *****
237
+ [2025-07-10 16:38:52,380][transformers.trainer][INFO] - Num examples = 132
238
+ [2025-07-10 16:38:52,380][transformers.trainer][INFO] - Batch size = 4
239
+ [2025-07-10 16:39:00,857][transformers.trainer][INFO] - Saving model checkpoint to /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-160
240
+ [2025-07-10 16:39:00,862][transformers.configuration_utils][INFO] - Configuration saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-160/config.json
241
+ [2025-07-10 16:39:10,224][transformers.modeling_utils][INFO] - Model weights saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-160/model.safetensors
242
+ [2025-07-10 16:46:16,187][transformers.trainer][INFO] - The following columns in the Evaluation set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
243
+ [2025-07-10 16:46:16,213][transformers.trainer][INFO] -
244
+ ***** Running Evaluation *****
245
+ [2025-07-10 16:46:16,214][transformers.trainer][INFO] - Num examples = 132
246
+ [2025-07-10 16:46:16,214][transformers.trainer][INFO] - Batch size = 4
247
+ [2025-07-10 16:46:24,732][transformers.trainer][INFO] - Saving model checkpoint to /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-192
248
+ [2025-07-10 16:46:24,736][transformers.configuration_utils][INFO] - Configuration saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-192/config.json
249
+ [2025-07-10 16:46:31,956][transformers.modeling_utils][INFO] - Model weights saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-192/model.safetensors
250
+ [2025-07-10 16:46:38,914][transformers.trainer][INFO] - Deleting older checkpoint [/workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-160] due to args.save_total_limit
251
+ [2025-07-10 16:53:35,545][transformers.trainer][INFO] - The following columns in the Evaluation set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
252
+ [2025-07-10 16:53:35,550][transformers.trainer][INFO] -
253
+ ***** Running Evaluation *****
254
+ [2025-07-10 16:53:35,550][transformers.trainer][INFO] - Num examples = 132
255
+ [2025-07-10 16:53:35,550][transformers.trainer][INFO] - Batch size = 4
256
+ [2025-07-10 16:53:44,013][transformers.trainer][INFO] - Saving model checkpoint to /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-224
257
+ [2025-07-10 16:53:44,015][transformers.configuration_utils][INFO] - Configuration saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-224/config.json
258
+ [2025-07-10 16:53:48,622][transformers.modeling_utils][INFO] - Model weights saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-224/model.safetensors
259
+ [2025-07-10 16:53:55,055][transformers.trainer][INFO] - Deleting older checkpoint [/workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-192] due to args.save_total_limit
260
+ [2025-07-10 17:00:51,991][transformers.trainer][INFO] - The following columns in the Evaluation set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
261
+ [2025-07-10 17:00:51,995][transformers.trainer][INFO] -
262
+ ***** Running Evaluation *****
263
+ [2025-07-10 17:00:51,995][transformers.trainer][INFO] - Num examples = 132
264
+ [2025-07-10 17:00:51,995][transformers.trainer][INFO] - Batch size = 4
265
+ [2025-07-10 17:01:00,444][transformers.trainer][INFO] - Saving model checkpoint to /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-256
266
+ [2025-07-10 17:01:00,447][transformers.configuration_utils][INFO] - Configuration saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-256/config.json
267
+ [2025-07-10 17:01:05,603][transformers.modeling_utils][INFO] - Model weights saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-256/model.safetensors
268
+ [2025-07-10 17:01:12,787][transformers.trainer][INFO] - Deleting older checkpoint [/workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-224] due to args.save_total_limit
269
+ [2025-07-10 17:08:09,449][transformers.trainer][INFO] - The following columns in the Evaluation set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
270
+ [2025-07-10 17:08:09,456][transformers.trainer][INFO] -
271
+ ***** Running Evaluation *****
272
+ [2025-07-10 17:08:09,456][transformers.trainer][INFO] - Num examples = 132
273
+ [2025-07-10 17:08:09,456][transformers.trainer][INFO] - Batch size = 4
274
+ [2025-07-10 17:08:17,939][transformers.trainer][INFO] - Saving model checkpoint to /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-288
275
+ [2025-07-10 17:08:17,941][transformers.configuration_utils][INFO] - Configuration saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-288/config.json
276
+ [2025-07-10 17:08:24,232][transformers.modeling_utils][INFO] - Model weights saved in /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-288/model.safetensors
277
+ [2025-07-10 17:08:33,143][transformers.trainer][INFO] - Deleting older checkpoint [/workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-256] due to args.save_total_limit
278
+ [2025-07-10 17:08:33,643][transformers.trainer][INFO] -
279
+
280
+ Training completed. Do not forget to share your model on huggingface.co/models =)
281
+
282
+
283
+ [2025-07-10 17:08:33,644][transformers.trainer][INFO] - Loading best model from /workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-128 (score: 0.6863342898134864).
284
+ [2025-07-10 17:08:36,030][transformers.trainer][INFO] - Deleting older checkpoint [/workspace/jbcs2025/outputs/2025-07-10/15-53-06/results/checkpoint-288] due to args.save_total_limit
285
+ [2025-07-10 17:08:36,692][transformers.trainer][INFO] - The following columns in the Evaluation set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
286
+ [2025-07-10 17:08:36,696][transformers.trainer][INFO] -
287
+ ***** Running Evaluation *****
288
+ [2025-07-10 17:08:36,696][transformers.trainer][INFO] - Num examples = 132
289
+ [2025-07-10 17:08:36,696][transformers.trainer][INFO] - Batch size = 4
290
+ [2025-07-10 17:08:45,160][__main__][INFO] - Training completed successfully.
291
+ [2025-07-10 17:08:45,160][__main__][INFO] - Running on Test
292
+ [2025-07-10 17:08:45,161][transformers.trainer][INFO] - The following columns in the Evaluation set don't have a corresponding argument in `DebertaV2ForSequenceClassification.forward` and have been ignored: prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades. If prompt, id, essay_year, id_prompt, supporting_text, reference, essay_text, grades are not expected by `DebertaV2ForSequenceClassification.forward`, you can safely ignore this message.
293
+ [2025-07-10 17:08:45,164][transformers.trainer][INFO] -
294
+ ***** Running Evaluation *****
295
+ [2025-07-10 17:08:45,164][transformers.trainer][INFO] - Num examples = 138
296
+ [2025-07-10 17:08:45,164][transformers.trainer][INFO] - Batch size = 4
297
+ [2025-07-10 17:08:53,928][__main__][INFO] - Test metrics: {'eval_loss': 0.9303792119026184, 'eval_model_preparation_time': 0.0074, 'eval_accuracy': 0.7028985507246377, 'eval_RMSE': 25.931906372573962, 'eval_QWK': 0.6826328310864394, 'eval_HDIV': 0.007246376811594235, 'eval_Macro_F1': 0.49688027796692447, 'eval_Micro_F1': 0.7028985507246377, 'eval_Weighted_F1': 0.7116672747290037, 'eval_TP_0': 0, 'eval_TN_0': 137, 'eval_FP_0': 0, 'eval_FN_0': 1, 'eval_TP_1': 0, 'eval_TN_1': 138, 'eval_FP_1': 0, 'eval_FN_1': 0, 'eval_TP_2': 6, 'eval_TN_2': 122, 'eval_FP_2': 6, 'eval_FN_2': 4, 'eval_TP_3': 45, 'eval_TN_3': 65, 'eval_FP_3': 7, 'eval_FN_3': 21, 'eval_TP_4': 40, 'eval_TN_4': 71, 'eval_FP_4': 16, 'eval_FN_4': 11, 'eval_TP_5': 6, 'eval_TN_5': 116, 'eval_FP_5': 12, 'eval_FN_5': 4, 'eval_runtime': 8.7537, 'eval_samples_per_second': 15.765, 'eval_steps_per_second': 3.998, 'epoch': 9.0}
298
+ [2025-07-10 17:08:53,929][transformers.trainer][INFO] - Saving model checkpoint to ./results/best_model
299
+ [2025-07-10 17:08:53,931][transformers.configuration_utils][INFO] - Configuration saved in ./results/best_model/config.json
300
+ [2025-07-10 17:08:58,786][transformers.modeling_utils][INFO] - Model weights saved in ./results/best_model/model.safetensors
301
+ [2025-07-10 17:08:58,789][transformers.tokenization_utils_base][INFO] - tokenizer config file saved in ./results/best_model/tokenizer_config.json
302
+ [2025-07-10 17:08:58,789][transformers.tokenization_utils_base][INFO] - Special tokens file saved in ./results/best_model/special_tokens_map.json
303
+ [2025-07-10 17:08:58,816][__main__][INFO] - Model and tokenizer saved to ./results/best_model
304
+ [2025-07-10 17:08:58,940][__main__][INFO] - Fine Tuning Finished.
305
+ [2025-07-10 17:08:59,455][__main__][INFO] - Total emissions: 0.1036 kg CO2eq
special_tokens_map.json ADDED
@@ -0,0 +1,51 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "bos_token": {
3
+ "content": "[CLS]",
4
+ "lstrip": false,
5
+ "normalized": false,
6
+ "rstrip": false,
7
+ "single_word": false
8
+ },
9
+ "cls_token": {
10
+ "content": "[CLS]",
11
+ "lstrip": false,
12
+ "normalized": false,
13
+ "rstrip": false,
14
+ "single_word": false
15
+ },
16
+ "eos_token": {
17
+ "content": "[SEP]",
18
+ "lstrip": false,
19
+ "normalized": false,
20
+ "rstrip": false,
21
+ "single_word": false
22
+ },
23
+ "mask_token": {
24
+ "content": "[MASK]",
25
+ "lstrip": false,
26
+ "normalized": false,
27
+ "rstrip": false,
28
+ "single_word": false
29
+ },
30
+ "pad_token": {
31
+ "content": "[PAD]",
32
+ "lstrip": false,
33
+ "normalized": false,
34
+ "rstrip": false,
35
+ "single_word": false
36
+ },
37
+ "sep_token": {
38
+ "content": "[SEP]",
39
+ "lstrip": false,
40
+ "normalized": false,
41
+ "rstrip": false,
42
+ "single_word": false
43
+ },
44
+ "unk_token": {
45
+ "content": "[UNK]",
46
+ "lstrip": false,
47
+ "normalized": true,
48
+ "rstrip": false,
49
+ "single_word": false
50
+ }
51
+ }
tokenizer.json ADDED
The diff for this file is too large to render. See raw diff
 
tokenizer_config.json ADDED
@@ -0,0 +1,59 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "added_tokens_decoder": {
3
+ "0": {
4
+ "content": "[PAD]",
5
+ "lstrip": false,
6
+ "normalized": false,
7
+ "rstrip": false,
8
+ "single_word": false,
9
+ "special": true
10
+ },
11
+ "1": {
12
+ "content": "[CLS]",
13
+ "lstrip": false,
14
+ "normalized": false,
15
+ "rstrip": false,
16
+ "single_word": false,
17
+ "special": true
18
+ },
19
+ "2": {
20
+ "content": "[SEP]",
21
+ "lstrip": false,
22
+ "normalized": false,
23
+ "rstrip": false,
24
+ "single_word": false,
25
+ "special": true
26
+ },
27
+ "3": {
28
+ "content": "[UNK]",
29
+ "lstrip": false,
30
+ "normalized": true,
31
+ "rstrip": false,
32
+ "single_word": false,
33
+ "special": true
34
+ },
35
+ "128000": {
36
+ "content": "[MASK]",
37
+ "lstrip": false,
38
+ "normalized": false,
39
+ "rstrip": false,
40
+ "single_word": false,
41
+ "special": true
42
+ }
43
+ },
44
+ "bos_token": "[CLS]",
45
+ "clean_up_tokenization_spaces": true,
46
+ "cls_token": "[CLS]",
47
+ "do_lower_case": false,
48
+ "eos_token": "[SEP]",
49
+ "extra_special_tokens": {},
50
+ "mask_token": "[MASK]",
51
+ "model_max_length": 512,
52
+ "pad_token": "[PAD]",
53
+ "sep_token": "[SEP]",
54
+ "sp_model_kwargs": {},
55
+ "split_by_punct": false,
56
+ "tokenizer_class": "DebertaV2Tokenizer",
57
+ "unk_token": "[UNK]",
58
+ "vocab_type": "spm"
59
+ }
training_args.bin ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:4cf94482fd02c284d5af1194f8d9537d3c5a6676656bebb88ad94465b29991aa
3
+ size 5777