Training in progress, step 2
Browse files- adapter_config.json +29 -0
- adapter_model.safetensors +3 -0
- runs/Nov28_14-57-04_gradientresourceone/events.out.tfevents.1732802226.gradientresourceone.154389.0 +3 -0
- runs/Nov28_15-00-39_gradientresourceone/events.out.tfevents.1732802440.gradientresourceone.155303.0 +3 -0
- runs/Nov28_15-08-08_gradientresourceone/events.out.tfevents.1732802888.gradientresourceone.157318.0 +3 -0
- runs/Nov28_15-14-02_gradientresourceone/events.out.tfevents.1732803243.gradientresourceone.158689.0 +3 -0
- runs/Nov28_15-16-02_gradientresourceone/events.out.tfevents.1732803363.gradientresourceone.159408.0 +3 -0
- runs/Nov28_15-20-26_gradientresourceone/events.out.tfevents.1732803627.gradientresourceone.160534.0 +3 -0
- runs/Nov28_15-24-47_gradientresourceone/events.out.tfevents.1732803888.gradientresourceone.161752.0 +3 -0
- runs/Nov28_15-46-02_gradientresourceone/events.out.tfevents.1732805163.gradientresourceone.166070.0 +3 -0
- runs/Nov28_16-10-52_gradientresourceone/events.out.tfevents.1732806653.gradientresourceone.171421.0 +3 -0
- runs/Nov28_16-17-54_gradientresourceone/events.out.tfevents.1732807075.gradientresourceone.173268.0 +3 -0
- runs/Nov28_16-22-17_gradientresourceone/events.out.tfevents.1732807338.gradientresourceone.174485.0 +3 -0
- special_tokens_map.json +24 -0
- tokenizer.model +3 -0
- tokenizer_config.json +0 -0
- training_args.bin +3 -0
adapter_config.json
ADDED
@@ -0,0 +1,29 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"alpha_pattern": {},
|
3 |
+
"auto_mapping": null,
|
4 |
+
"base_model_name_or_path": "mistralai/Mistral-7B-Instruct-v0.3",
|
5 |
+
"bias": "none",
|
6 |
+
"fan_in_fan_out": false,
|
7 |
+
"inference_mode": true,
|
8 |
+
"init_lora_weights": true,
|
9 |
+
"layer_replication": null,
|
10 |
+
"layers_pattern": null,
|
11 |
+
"layers_to_transform": null,
|
12 |
+
"loftq_config": {},
|
13 |
+
"lora_alpha": 16,
|
14 |
+
"lora_dropout": 0.1,
|
15 |
+
"megatron_config": null,
|
16 |
+
"megatron_core": "megatron.core",
|
17 |
+
"modules_to_save": null,
|
18 |
+
"peft_type": "LORA",
|
19 |
+
"r": 4,
|
20 |
+
"rank_pattern": {},
|
21 |
+
"revision": null,
|
22 |
+
"target_modules": [
|
23 |
+
"v_proj",
|
24 |
+
"q_proj"
|
25 |
+
],
|
26 |
+
"task_type": "CAUSAL_LM",
|
27 |
+
"use_dora": false,
|
28 |
+
"use_rslora": false
|
29 |
+
}
|
adapter_model.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:f40a6a4b0ee20dd8db6b19e6c0f00270f94825ef0f411309c3b75e1b86c2f8f5
|
3 |
+
size 6832600
|
runs/Nov28_14-57-04_gradientresourceone/events.out.tfevents.1732802226.gradientresourceone.154389.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:509c7924d6d84356dbe3b185adefcd5e6ca43d0b2d51a0b104da3e1f16f7af18
|
3 |
+
size 5110
|
runs/Nov28_15-00-39_gradientresourceone/events.out.tfevents.1732802440.gradientresourceone.155303.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:7403ba401a5976a3744d5dff99768e05eb4667a2322e4a72abab6342dfde00d7
|
3 |
+
size 5731
|
runs/Nov28_15-08-08_gradientresourceone/events.out.tfevents.1732802888.gradientresourceone.157318.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:b0f488cef07176722634741fa3b466b934e5f9c9e2071363ed0bcc847c667598
|
3 |
+
size 5591
|
runs/Nov28_15-14-02_gradientresourceone/events.out.tfevents.1732803243.gradientresourceone.158689.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:b9e46dff0c4954bb8286edf7ad43f8a9ff478782124d7bce14f4f3d8290f371f
|
3 |
+
size 5591
|
runs/Nov28_15-16-02_gradientresourceone/events.out.tfevents.1732803363.gradientresourceone.159408.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:6edabc82b9865a65f27c8a8afb4be171c1456137d78aeaeebbc2b41214b08ef7
|
3 |
+
size 5591
|
runs/Nov28_15-20-26_gradientresourceone/events.out.tfevents.1732803627.gradientresourceone.160534.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:ebf29fe58b5c749785dd072affcf7482f081294e7d8c223617e7f3ca1954359e
|
3 |
+
size 5524
|
runs/Nov28_15-24-47_gradientresourceone/events.out.tfevents.1732803888.gradientresourceone.161752.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:103e9c590481ec485718347207aa57aac954bb094d85f38f777bcd5b9fdd9e95
|
3 |
+
size 5593
|
runs/Nov28_15-46-02_gradientresourceone/events.out.tfevents.1732805163.gradientresourceone.166070.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:a49e07564e34be4fa2c6175dadce36f409fa1a8fec3bbf590f21f2e215755e0f
|
3 |
+
size 6004
|
runs/Nov28_16-10-52_gradientresourceone/events.out.tfevents.1732806653.gradientresourceone.171421.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:fee267d9c547dc9e0c6025135743128a16a72b606457bae6327f7c39c48f3199
|
3 |
+
size 6558
|
runs/Nov28_16-17-54_gradientresourceone/events.out.tfevents.1732807075.gradientresourceone.173268.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:cbb6bde008e67832921c3f908f8002ff7a54e63e8bf9db46ca1fc5d731f49415
|
3 |
+
size 5937
|
runs/Nov28_16-22-17_gradientresourceone/events.out.tfevents.1732807338.gradientresourceone.174485.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:739fc6d81358a35f2714bb4b8faf31b0aff6ed0d58a0c3c51b4de48e545573cf
|
3 |
+
size 5872
|
special_tokens_map.json
ADDED
@@ -0,0 +1,24 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"bos_token": {
|
3 |
+
"content": "<s>",
|
4 |
+
"lstrip": false,
|
5 |
+
"normalized": false,
|
6 |
+
"rstrip": false,
|
7 |
+
"single_word": false
|
8 |
+
},
|
9 |
+
"eos_token": {
|
10 |
+
"content": "</s>",
|
11 |
+
"lstrip": false,
|
12 |
+
"normalized": false,
|
13 |
+
"rstrip": false,
|
14 |
+
"single_word": false
|
15 |
+
},
|
16 |
+
"pad_token": "</s>",
|
17 |
+
"unk_token": {
|
18 |
+
"content": "<unk>",
|
19 |
+
"lstrip": false,
|
20 |
+
"normalized": false,
|
21 |
+
"rstrip": false,
|
22 |
+
"single_word": false
|
23 |
+
}
|
24 |
+
}
|
tokenizer.model
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:37f00374dea48658ee8f5d0f21895b9bc55cb0103939607c8185bfd1c6ca1f89
|
3 |
+
size 587404
|
tokenizer_config.json
ADDED
The diff for this file is too large to render.
See raw diff
|
|
training_args.bin
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:ff1cfea0e866632c1d9da6fac625c510f108a5ba592f98733d839267d5c006ee
|
3 |
+
size 5304
|