### model
model_name_or_path: /MODEL_PATH
trust_remote_code: true


### method
stage: sft
do_train: true
finetuning_type: full
freeze_vision_tower: false
freeze_multi_modal_projector: false
deepspeed: /DEEPSPEED_CONFIG_PATH

### dataset
dataset: rft_train_llamafactory_dt0
dataset_dir: /DATASET_DIR
template: qwen3_vl_nothink
cutoff_len: 8192
overwrite_cache: false
preprocessing_num_workers: 256

### output
output_dir: /OUTPUT_DIR
logging_steps: 10
save_steps: 5000
plot_loss: true
overwrite_output_dir: true

### train
per_device_train_batch_size: 4
gradient_accumulation_steps: 8
learning_rate: 1.0e-5
num_train_epochs: 3.0
lr_scheduler_type: cosine
warmup_ratio: 0.1
bf16: true
ddp_timeout: 180000000

### eval
val_size: 0.1
per_device_eval_batch_size: 2
eval_strategy: steps
eval_steps: 45

