Spaces:

Dovakiins
/

qwerrwe

Build error

App Files Files Community

Nanobit commited on Aug 12, 2023

Commit

a276c9c

unverified ·

1 Parent(s): 7019509

Fix(save): Save as safetensors (#363)

Browse files

Files changed (1) hide show

scripts/finetune.py +10 -5

scripts/finetune.py CHANGED Viewed

@@ -257,6 +257,8 @@ def train(
     LOG.info("loading model and peft_config...")
     model, peft_config = load_model(cfg, tokenizer)
     if "merge_lora" in kwargs and cfg.adapter is not None:
         LOG.info("running merge of LoRA with base model")
         model = model.merge_and_unload()
@@ -264,7 +266,10 @@ def train(
         if cfg.local_rank == 0:
             LOG.info("saving merged model")
-            model.save_pretrained(str(Path(cfg.output_dir) / "merged"))
             tokenizer.save_pretrained(str(Path(cfg.output_dir) / "merged"))
         return
@@ -280,7 +285,7 @@ def train(
         return
     if "shard" in kwargs:
-        model.save_pretrained(cfg.output_dir)
         return
     trainer = setup_trainer(cfg, train_dataset, eval_dataset, model, tokenizer)
@@ -302,7 +307,7 @@ def train(
         def terminate_handler(_, __, model):
             if cfg.flash_optimum:
                 model = BetterTransformer.reverse(model)
-            model.save_pretrained(cfg.output_dir)
             sys.exit(0)
         signal.signal(
@@ -342,11 +347,11 @@ def train(
     # TODO do we need this fix? https://huggingface.co/docs/accelerate/usage_guides/fsdp#saving-and-loading
     # only save on rank 0, otherwise it corrupts output on multi-GPU when multiple processes attempt to write the same file
     if cfg.fsdp:
-        model.save_pretrained(cfg.output_dir)
     elif cfg.local_rank == 0:
         if cfg.flash_optimum:
             model = BetterTransformer.reverse(model)
-        model.save_pretrained(cfg.output_dir)
     # trainer.save_model(cfg.output_dir)  # TODO this may be needed for deepspeed to work? need to review another time

     LOG.info("loading model and peft_config...")
     model, peft_config = load_model(cfg, tokenizer)
+    safe_serialization = cfg.save_safetensors is True
     if "merge_lora" in kwargs and cfg.adapter is not None:
         LOG.info("running merge of LoRA with base model")
         model = model.merge_and_unload()
         if cfg.local_rank == 0:
             LOG.info("saving merged model")
+            model.save_pretrained(
+                str(Path(cfg.output_dir) / "merged"),
+                safe_serialization=safe_serialization,
+            )
             tokenizer.save_pretrained(str(Path(cfg.output_dir) / "merged"))
         return
         return
     if "shard" in kwargs:
+        model.save_pretrained(cfg.output_dir, safe_serialization=safe_serialization)
         return
     trainer = setup_trainer(cfg, train_dataset, eval_dataset, model, tokenizer)
         def terminate_handler(_, __, model):
             if cfg.flash_optimum:
                 model = BetterTransformer.reverse(model)
+            model.save_pretrained(cfg.output_dir, safe_serialization=safe_serialization)
             sys.exit(0)
         signal.signal(
     # TODO do we need this fix? https://huggingface.co/docs/accelerate/usage_guides/fsdp#saving-and-loading
     # only save on rank 0, otherwise it corrupts output on multi-GPU when multiple processes attempt to write the same file
     if cfg.fsdp:
+        model.save_pretrained(cfg.output_dir, safe_serialization=safe_serialization)
     elif cfg.local_rank == 0:
         if cfg.flash_optimum:
             model = BetterTransformer.reverse(model)
+        model.save_pretrained(cfg.output_dir, safe_serialization=safe_serialization)
     # trainer.save_model(cfg.output_dir)  # TODO this may be needed for deepspeed to work? need to review another time