fix conditional for None values

handle batch size correchtly when using split and dispatch batches
2025-08-17 12:49:48 -04:00 · 2025-08-16 22:05:31 -04:00
3 changed files with 15 additions and 1 deletions
--- a/src/axolotl/core/builders/causal.py
+++ b/src/axolotl/core/builders/causal.py
@@ -424,7 +424,7 @@ class HFCausalTrainerBuilder(TrainerBuilderBase):
    ):
        if training_args.pretraining:
            if (
-                self.cfg.pretraining_sample_concatenation is False
+                not self.cfg.pretraining_sample_concatenation
                or self.cfg.micro_batch_size > 1
            ):
                return DataCollatorForSeq2Seq(self.tokenizer, **kwargs)
--- a/src/axolotl/core/trainers/base.py
+++ b/src/axolotl/core/trainers/base.py
@@ -272,6 +272,20 @@ class AxolotlTrainer(
                    num_workers=self.args.dataloader_num_workers,
                    rank=self.args.process_index,
                )
+
+        if (
+            self.args.accelerator_config is not None
+            and self.args.accelerator_config.split_batches
+            and self.args.accelerator_config.dispatch_batches
+        ):
+            if self.args.sample_packing and self.args.pretraining:
+                if not self.args.eval_sample_packing and not is_training:
+                    dataloader_params["batch_size"] *= self.accelerator.num_processes
+                else:
+                    dataloader_params["batch_size"] = self.accelerator.num_processes
+            elif not self.args.sample_packing and self.args.pretraining:
+                dataloader_params["batch_size"] *= self.accelerator.num_processes
+
        if self.args.sample_packing and (
            (is_training and not self.args.pretraining)
            or (not is_training and self.args.eval_sample_packing is not False)
--- a/src/axolotl/exception_handling.py
+++ b/src/axolotl/exception_handling.py
Author	SHA1	Message	Date
Wing Lian	bb65157dcf	fix conditional for None values	2025-08-17 12:49:48 -04:00
Wing Lian	7fd3d8abc4	handle batch size correchtly when using split and dispatch batches	2025-08-16 22:05:31 -04:00