{ "alpha_pattern": {}, "auto_mapping": { "base_model_class": "LLaDAModelLM", "parent_library": "dllm.pipelines.llada.models.modeling_llada" }, "base_model_name_or_path": "GSAI-ML/LLaDA-8B-Instruct", "bias": "none", "corda_config": null, "eva_config": null, "exclude_modules": null, "fan_in_fan_out": false, "inference_mode": true, "init_lora_weights": true, "layer_replication": null, "layers_pattern": null, "layers_to_transform": null, "loftq_config": {}, "lora_alpha": 64, "lora_bias": false, "lora_dropout": 0.0, "megatron_config": null, "megatron_core": "megatron.core", "modules_to_save": null, "peft_type": "LORA", "qalora_group_size": 16, "r": 32, "rank_pattern": {}, "revision": null, "target_modules": [ "21.ff_out", "19.ff_out", "28.ff_out", "2.ff_out", "14.ff_out", "7.ff_out", "30.ff_out", "4.ff_out", "26.ff_out", "15.ff_out", "k_proj", "25.ff_out", "16.ff_out", "9.ff_out", "20.ff_out", "3.ff_out", "18.ff_out", "11.ff_out", "0.ff_out", "13.ff_out", "29.ff_out", "ff_proj", "17.ff_out", "31.ff_out", "23.ff_out", "22.ff_out", "12.ff_out", "8.ff_out", "24.ff_out", "5.ff_out", "27.ff_out", "v_proj", "6.ff_out", "1.ff_out", "10.ff_out", "q_proj", "up_proj", "attn_out" ], "target_parameters": null, "task_type": "CAUSAL_LM", "trainable_token_indices": null, "use_dora": false, "use_qalora": false, "use_rslora": false }