Training in progress, step 125

Files changed (4) hide show

README.md CHANGED Viewed

@@ -4,8 +4,8 @@ library_name: transformers
 model_name: mistral-supervised
 tags:
 - generated_from_trainer
-- trl
 - sft
 licence: license
 ---
@@ -27,13 +27,14 @@ print(output["generated_text"])
 ## Training procedure
 This model was trained with SFT.
 ### Framework versions
-- TRL: 0.12.0
 - Transformers: 4.54.0
 - Pytorch: 2.7.1
 - Datasets: 4.0.0
@@ -48,7 +49,7 @@ Cite TRL as:
 ```bibtex
 @misc{vonwerra2022trl,
 	title        = {{TRL: Transformer Reinforcement Learning}},
-	author       = {Leandro von Werra and Younes Belkada and Lewis Tunstall and Edward Beeching and Tristan Thrush and Nathan Lambert and Shengyi Huang and Kashif Rasul and Quentin Gallouédec},
 	year         = 2020,
 	journal      = {GitHub repository},
 	publisher    = {GitHub},

 model_name: mistral-supervised
 tags:
 - generated_from_trainer
 - sft
+- trl
 licence: license
 ---
 ## Training procedure
 This model was trained with SFT.
 ### Framework versions
+- TRL: 0.20.0
 - Transformers: 4.54.0
 - Pytorch: 2.7.1
 - Datasets: 4.0.0
 ```bibtex
 @misc{vonwerra2022trl,
 	title        = {{TRL: Transformer Reinforcement Learning}},
+	author       = {Leandro von Werra and Younes Belkada and Lewis Tunstall and Edward Beeching and Tristan Thrush and Nathan Lambert and Shengyi Huang and Kashif Rasul and Quentin Gallou{\'e}dec},
 	year         = 2020,
 	journal      = {GitHub repository},
 	publisher    = {GitHub},

model.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:98f5ae2aad1cdd22c6ee19f0df9ef9f79f4537e1d22d80056d9474550072f877
 size 338197712

 version https://git-lfs.github.com/spec/v1
+oid sha256:0299c3d0a1274718f7072c15573620fe605875f8d06e2c221d0bffe67685eb19
 size 338197712

tokenizer.json CHANGED Viewed

@@ -1,11 +1,6 @@
 {
   "version": "1.0",
-  "truncation": {
-    "direction": "Right",
-    "max_length": 512,
-    "strategy": "LongestFirst",
-    "stride": 0
-  },
   "padding": null,
   "added_tokens": [
     {

 {
   "version": "1.0",
+  "truncation": null,
   "padding": null,
   "added_tokens": [
     {

training_args.bin CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:b2229bae4bb325b571e23eafc16b133a7b62c79e3f9994ca8a140e5ed1efc65a
-size 5969

 version https://git-lfs.github.com/spec/v1
+oid sha256:5a4cbc2c8f0f687165bfcbcc5879bac2fe8161afd799e83f593608ff9a333d41
+size 6097