Text Classification
Transformers
TensorBoard
Safetensors
bert
glue
rte
max_length_128
dropout_0.4
Generated from Trainer
text-embeddings-inference
Instructions to use ipeksnmz/bert-base-uncased-finetuned-rte-run_3 with libraries, inference providers, notebooks, and local apps. Follow these links to get started.
- Libraries
- Transformers
How to use ipeksnmz/bert-base-uncased-finetuned-rte-run_3 with Transformers:
# Use a pipeline as a high-level helper from transformers import pipeline pipe = pipeline("text-classification", model="ipeksnmz/bert-base-uncased-finetuned-rte-run_3")# Load model directly from transformers import AutoTokenizer, AutoModelForSequenceClassification tokenizer = AutoTokenizer.from_pretrained("ipeksnmz/bert-base-uncased-finetuned-rte-run_3") model = AutoModelForSequenceClassification.from_pretrained("ipeksnmz/bert-base-uncased-finetuned-rte-run_3", device_map="auto") - Notebooks
- Google Colab
- Kaggle
Training in progress, epoch 1
Browse files- model.safetensors +1 -1
- run-3/checkpoint-78/model.safetensors +1 -1
- run-3/checkpoint-78/optimizer.pt +1 -1
- run-3/checkpoint-78/scheduler.pt +1 -1
- run-3/checkpoint-78/trainer_state.json +11 -11
- run-3/checkpoint-78/training_args.bin +1 -1
- runs/Apr07_19-56-35_578d22db0b7b/events.out.tfevents.1744056356.578d22db0b7b.632.5 +3 -0
- training_args.bin +1 -1
model.safetensors
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 437958648
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:cc777af6911bc2f91e0cdceced703b8e4a09d550e56d6e621ac9c0b57def38f8
|
| 3 |
size 437958648
|
run-3/checkpoint-78/model.safetensors
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 437958648
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:cc777af6911bc2f91e0cdceced703b8e4a09d550e56d6e621ac9c0b57def38f8
|
| 3 |
size 437958648
|
run-3/checkpoint-78/optimizer.pt
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 876038394
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:93e4a5484799e0792b95d92b1eec17806230e691d692a57a2353e9e3966cfffa
|
| 3 |
size 876038394
|
run-3/checkpoint-78/scheduler.pt
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 1064
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:ae41ec8a4c453ec1e587897e4b3ba19337bd7deccb146115ffdc1040dee239e0
|
| 3 |
size 1064
|
run-3/checkpoint-78/trainer_state.json
CHANGED
|
@@ -1,6 +1,6 @@
|
|
| 1 |
{
|
| 2 |
"best_global_step": 78,
|
| 3 |
-
"best_metric": 0.
|
| 4 |
"best_model_checkpoint": "bert-base-uncased-finetuned-rte-run_3/run-3/checkpoint-78",
|
| 5 |
"epoch": 1.0,
|
| 6 |
"eval_steps": 500,
|
|
@@ -11,18 +11,18 @@
|
|
| 11 |
"log_history": [
|
| 12 |
{
|
| 13 |
"epoch": 1.0,
|
| 14 |
-
"eval_accuracy": 0.
|
| 15 |
-
"eval_loss": 0.
|
| 16 |
-
"eval_runtime": 0.
|
| 17 |
-
"eval_samples_per_second":
|
| 18 |
-
"eval_steps_per_second": 19.
|
| 19 |
"step": 78
|
| 20 |
}
|
| 21 |
],
|
| 22 |
"logging_steps": 500,
|
| 23 |
-
"max_steps":
|
| 24 |
"num_input_tokens_seen": 0,
|
| 25 |
-
"num_train_epochs":
|
| 26 |
"save_steps": 500,
|
| 27 |
"stateful_callbacks": {
|
| 28 |
"TrainerControl": {
|
|
@@ -31,7 +31,7 @@
|
|
| 31 |
"should_evaluate": false,
|
| 32 |
"should_log": false,
|
| 33 |
"should_save": true,
|
| 34 |
-
"should_training_stop":
|
| 35 |
},
|
| 36 |
"attributes": {}
|
| 37 |
}
|
|
@@ -41,9 +41,9 @@
|
|
| 41 |
"trial_name": null,
|
| 42 |
"trial_params": {
|
| 43 |
"dropout": 0.1,
|
| 44 |
-
"learning_rate":
|
| 45 |
"max_length": 256,
|
| 46 |
-
"num_train_epochs":
|
| 47 |
"per_device_train_batch_size": 32
|
| 48 |
}
|
| 49 |
}
|
|
|
|
| 1 |
{
|
| 2 |
"best_global_step": 78,
|
| 3 |
+
"best_metric": 0.4693140794223827,
|
| 4 |
"best_model_checkpoint": "bert-base-uncased-finetuned-rte-run_3/run-3/checkpoint-78",
|
| 5 |
"epoch": 1.0,
|
| 6 |
"eval_steps": 500,
|
|
|
|
| 11 |
"log_history": [
|
| 12 |
{
|
| 13 |
"epoch": 1.0,
|
| 14 |
+
"eval_accuracy": 0.4693140794223827,
|
| 15 |
+
"eval_loss": 0.6942658424377441,
|
| 16 |
+
"eval_runtime": 0.4698,
|
| 17 |
+
"eval_samples_per_second": 589.581,
|
| 18 |
+
"eval_steps_per_second": 19.156,
|
| 19 |
"step": 78
|
| 20 |
}
|
| 21 |
],
|
| 22 |
"logging_steps": 500,
|
| 23 |
+
"max_steps": 78,
|
| 24 |
"num_input_tokens_seen": 0,
|
| 25 |
+
"num_train_epochs": 1,
|
| 26 |
"save_steps": 500,
|
| 27 |
"stateful_callbacks": {
|
| 28 |
"TrainerControl": {
|
|
|
|
| 31 |
"should_evaluate": false,
|
| 32 |
"should_log": false,
|
| 33 |
"should_save": true,
|
| 34 |
+
"should_training_stop": true
|
| 35 |
},
|
| 36 |
"attributes": {}
|
| 37 |
}
|
|
|
|
| 41 |
"trial_name": null,
|
| 42 |
"trial_params": {
|
| 43 |
"dropout": 0.1,
|
| 44 |
+
"learning_rate": 8.77558697928249e-06,
|
| 45 |
"max_length": 256,
|
| 46 |
+
"num_train_epochs": 1,
|
| 47 |
"per_device_train_batch_size": 32
|
| 48 |
}
|
| 49 |
}
|
run-3/checkpoint-78/training_args.bin
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 5432
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:b6e3245604cac2709b0dab5c3ce10afa8e9bee1f4e19b928623f07264875b418
|
| 3 |
size 5432
|
runs/Apr07_19-56-35_578d22db0b7b/events.out.tfevents.1744056356.578d22db0b7b.632.5
ADDED
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:410d23ebace6e1b3b9ba889894b314b834382a58d7ec472feabfa7d2d798a265
|
| 3 |
+
size 5765
|
training_args.bin
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 5432
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:b6e3245604cac2709b0dab5c3ce10afa8e9bee1f4e19b928623f07264875b418
|
| 3 |
size 5432
|