fzzhang commited on
Commit
1546487
·
verified ·
1 Parent(s): a82ac7f

End of training

Browse files
README.md CHANGED
@@ -33,7 +33,7 @@ More information needed
33
  ### Training hyperparameters
34
 
35
  The following hyperparameters were used during training:
36
- - learning_rate: 2.5e-05
37
  - train_batch_size: 4
38
  - eval_batch_size: 8
39
  - seed: 0
 
33
  ### Training hyperparameters
34
 
35
  The following hyperparameters were used during training:
36
+ - learning_rate: 1e-05
37
  - train_batch_size: 4
38
  - eval_batch_size: 8
39
  - seed: 0
adapter_config.json CHANGED
@@ -22,9 +22,9 @@
22
  "spectral_top": true,
23
  "target_modules": [
24
  "o_proj",
25
- "gate_proj",
26
  "k_proj",
27
  "v_proj",
 
28
  "q_proj"
29
  ],
30
  "task_type": "CAUSAL_LM",
 
22
  "spectral_top": true,
23
  "target_modules": [
24
  "o_proj",
 
25
  "k_proj",
26
  "v_proj",
27
+ "gate_proj",
28
  "q_proj"
29
  ],
30
  "task_type": "CAUSAL_LM",
adapter_model.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:8f2369637f091818eee9dd8e4e4ac5f63d3d75a8836e446ee237fe35219b949d
3
  size 23111352
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:753cf432421eee1ace5987dcf430d8002e59e29c92450a17da78e6e90d72bb94
3
  size 23111352
runs/May19_11-17-47_mert-lambda-scalar/events.out.tfevents.1716142667.mert-lambda-scalar.4142014.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:2e26a5f884be85808cf25b19756f8ec445c2f7c1f1e45abd71591265cfb3638f
3
+ size 201994
training_args.bin CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:67e42acc01351cb172d8f7f31c9c618ab2e9c962f2916f8da3778b0d76b421f7
3
  size 4920
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:f48ded3f569be922bb9f0558d5cf084654a66b4e5a8aeb152d5417e8559c2235
3
  size 4920