checkpoint-522
Browse files- edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/README.md +0 -0
- edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/adapter_config.json +0 -0
- edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/adapter_model.safetensors +1 -1
- edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/optimizer.pt +1 -1
- edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/rng_state.pth +1 -1
- edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/scheduler.pt +1 -1
- edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/trainer_state.json +26 -4
- edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/training_args.bin +0 -0
- edit-qwen3.5-2b/tb/events.out.tfevents.1790604388.9b7939721266.4215.0 +2 -2
edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/README.md
RENAMED
|
File without changes
|
edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/adapter_config.json
RENAMED
|
File without changes
|
edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/adapter_model.safetensors
RENAMED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 134604152
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:5a131fb3173093e470e0ebf104035a1d8036da3476028965ffe8b4c29716abf5
|
| 3 |
size 134604152
|
edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/optimizer.pt
RENAMED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 269426431
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:16339f7d39dee8128ad4dad91ff053a817cadd0cd5175656d00073f85289b484
|
| 3 |
size 269426431
|
edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/rng_state.pth
RENAMED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 14645
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:e8da8f9385b4aef0617d60ce057ba8b334a00b95f8ef53b01233cb6867236f2f
|
| 3 |
size 14645
|
edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/scheduler.pt
RENAMED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 1465
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:57b1342a92a9abdc177d2ce815dae6b22c2b1803855905a8825da0b35bba438f
|
| 3 |
size 1465
|
edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/trainer_state.json
RENAMED
|
@@ -2,9 +2,9 @@
|
|
| 2 |
"best_global_step": 150,
|
| 3 |
"best_metric": 0.9355365037918091,
|
| 4 |
"best_model_checkpoint": "./outputs/edit-qwen3.5-2b/checkpoint-150",
|
| 5 |
-
"epoch":
|
| 6 |
"eval_steps": 50,
|
| 7 |
-
"global_step":
|
| 8 |
"is_hyper_param_search": false,
|
| 9 |
"is_local_process_zero": true,
|
| 10 |
"is_world_process_zero": true,
|
|
@@ -438,6 +438,28 @@
|
|
| 438 |
"eval_samples_per_second": 56.413,
|
| 439 |
"eval_steps_per_second": 14.103,
|
| 440 |
"step": 500
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 441 |
}
|
| 442 |
],
|
| 443 |
"logging_steps": 10,
|
|
@@ -452,12 +474,12 @@
|
|
| 452 |
"should_evaluate": false,
|
| 453 |
"should_log": false,
|
| 454 |
"should_save": true,
|
| 455 |
-
"should_training_stop":
|
| 456 |
},
|
| 457 |
"attributes": {}
|
| 458 |
}
|
| 459 |
},
|
| 460 |
-
"total_flos": 3.
|
| 461 |
"train_batch_size": 4,
|
| 462 |
"trial_name": null,
|
| 463 |
"trial_params": null
|
|
|
|
| 2 |
"best_global_step": 150,
|
| 3 |
"best_metric": 0.9355365037918091,
|
| 4 |
"best_model_checkpoint": "./outputs/edit-qwen3.5-2b/checkpoint-150",
|
| 5 |
+
"epoch": 3.0,
|
| 6 |
"eval_steps": 50,
|
| 7 |
+
"global_step": 522,
|
| 8 |
"is_hyper_param_search": false,
|
| 9 |
"is_local_process_zero": true,
|
| 10 |
"is_world_process_zero": true,
|
|
|
|
| 438 |
"eval_samples_per_second": 56.413,
|
| 439 |
"eval_steps_per_second": 14.103,
|
| 440 |
"step": 500
|
| 441 |
+
},
|
| 442 |
+
{
|
| 443 |
+
"epoch": 2.9337175792507204,
|
| 444 |
+
"grad_norm": 5.865486145019531,
|
| 445 |
+
"learning_rate": 2.6262626262626263e-06,
|
| 446 |
+
"loss": 0.3970566034317017,
|
| 447 |
+
"step": 510
|
| 448 |
+
},
|
| 449 |
+
{
|
| 450 |
+
"epoch": 2.9913544668587897,
|
| 451 |
+
"grad_norm": 0.623012363910675,
|
| 452 |
+
"learning_rate": 6.060606060606061e-07,
|
| 453 |
+
"loss": 0.2541007995605469,
|
| 454 |
+
"step": 520
|
| 455 |
+
},
|
| 456 |
+
{
|
| 457 |
+
"epoch": 3.0,
|
| 458 |
+
"eval_loss": 0.9820418953895569,
|
| 459 |
+
"eval_runtime": 2.547,
|
| 460 |
+
"eval_samples_per_second": 56.538,
|
| 461 |
+
"eval_steps_per_second": 14.134,
|
| 462 |
+
"step": 522
|
| 463 |
}
|
| 464 |
],
|
| 465 |
"logging_steps": 10,
|
|
|
|
| 474 |
"should_evaluate": false,
|
| 475 |
"should_log": false,
|
| 476 |
"should_save": true,
|
| 477 |
+
"should_training_stop": true
|
| 478 |
},
|
| 479 |
"attributes": {}
|
| 480 |
}
|
| 481 |
},
|
| 482 |
+
"total_flos": 3.448892658042317e+16,
|
| 483 |
"train_batch_size": 4,
|
| 484 |
"trial_name": null,
|
| 485 |
"trial_params": null
|
edit-qwen3.5-2b/{checkpoint-500 β checkpoint-522}/training_args.bin
RENAMED
|
File without changes
|
edit-qwen3.5-2b/tb/events.out.tfevents.1790604388.9b7939721266.4215.0
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
-
size
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:ef4878f89b7d5d077edeb7c9bce5f060004e2d63825b1031c34f943ff92d62d0
|
| 3 |
+
size 19324
|