HBboy commited on
Commit
0f8c413
·
verified ·
1 Parent(s): a47ca8d

Upload llamaboard_config.yaml with huggingface_hub

Browse files
Files changed (1) hide show
  1. llamaboard_config.yaml +78 -0
llamaboard_config.yaml ADDED
@@ -0,0 +1,78 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ top.booster: auto
2
+ top.checkpoint_path: null
3
+ top.finetuning_type: full
4
+ top.model_name: Qwen2.5-0.5B-Instruct
5
+ top.quantization_bit: none
6
+ top.quantization_method: bitsandbytes
7
+ top.rope_scaling: none
8
+ top.template: qwen
9
+ train.additional_target: ''
10
+ train.apollo_rank: 16
11
+ train.apollo_scale: 32
12
+ train.apollo_target: all
13
+ train.apollo_update_interval: 200
14
+ train.badam_mode: layer
15
+ train.badam_switch_interval: 50
16
+ train.badam_switch_mode: ascending
17
+ train.badam_update_ratio: 0.05
18
+ train.batch_size: 2
19
+ train.compute_type: bf16
20
+ train.create_new_adapter: false
21
+ train.cutoff_len: 2048
22
+ train.dataset:
23
+ - xiaosui-train
24
+ - identity
25
+ train.dataset_dir: data
26
+ train.ds_offload: false
27
+ train.ds_stage: none
28
+ train.extra_args: '{"optim": "adamw_torch"}'
29
+ train.freeze_extra_modules: ''
30
+ train.freeze_trainable_layers: 2
31
+ train.freeze_trainable_modules: all
32
+ train.galore_rank: 16
33
+ train.galore_scale: 2
34
+ train.galore_target: all
35
+ train.galore_update_interval: 200
36
+ train.gradient_accumulation_steps: 8
37
+ train.learning_rate: 5e-5
38
+ train.logging_steps: 5
39
+ train.lora_alpha: 16
40
+ train.lora_dropout: 0
41
+ train.lora_rank: 8
42
+ train.lora_target: ''
43
+ train.loraplus_lr_ratio: 0
44
+ train.lr_scheduler_type: cosine
45
+ train.mask_history: false
46
+ train.max_grad_norm: '1.0'
47
+ train.max_samples: '5000'
48
+ train.neat_packing: false
49
+ train.neftune_alpha: 0
50
+ train.num_train_epochs: '3.0'
51
+ train.packing: false
52
+ train.ppo_score_norm: false
53
+ train.ppo_whiten_rewards: false
54
+ train.pref_beta: 0.1
55
+ train.pref_ftx: 0
56
+ train.pref_loss: sigmoid
57
+ train.report_to:
58
+ - wandb
59
+ train.resize_vocab: false
60
+ train.reward_model: []
61
+ train.save_steps: 100
62
+ train.swanlab_api_key: ''
63
+ train.swanlab_mode: cloud
64
+ train.swanlab_project: llamafactory
65
+ train.swanlab_run_name: ''
66
+ train.swanlab_workspace: ''
67
+ train.train_on_prompt: false
68
+ train.training_stage: Supervised Fine-Tuning
69
+ train.use_apollo: false
70
+ train.use_badam: false
71
+ train.use_dora: false
72
+ train.use_galore: false
73
+ train.use_llama_pro: false
74
+ train.use_pissa: false
75
+ train.use_rslora: false
76
+ train.use_swanlab: false
77
+ train.val_size: 0
78
+ train.warmup_steps: 4