Completion only fine-tuning of instruction models with collections of HF datasets (#1103)

- Optional completion only fine-tuning with `--mask-prompt` - Collections of Hugging Face datasets --------- Co-authored-by: Awni Hannun <awni@apple.com>
2025-12-16 02:08:55 +08:00 · 2025-02-09 23:12:34 -05:00
parent 1ced1b00ca
commit 5865899c81
6 changed files with 199 additions and 85 deletions
--- a/llms/mlx_lm/lora.py
+++ b/llms/mlx_lm/lora.py
@@ -94,6 +94,14 @@ def build_parser():
        choices=["lora", "dora", "full"],
        help="Type of fine-tuning to perform: lora, dora, or full.",
    )
+
+    parser.add_argument(
+        "--mask-prompt",
+        action="store_true",
+        help="Mask the prompt in the loss when training",
+        default=False,
+    )
+
    parser.add_argument(
        "--num-layers",
        type=int,
@@ -219,6 +227,7 @@ def train_model(
            build_schedule(args.lr_schedule) if args.lr_schedule else args.learning_rate
        )
    )
+
    # Train model
    train(
        model=model,