fix: document offload gradient_checkpointing option (#2475)

2025-04-02 20:35:42 +07:00
parent a0117c9bce
commit adb593abac
1 changed files with 2 additions and 1 deletions
--- a/docs/config.qmd
+++ b/docs/config.qmd
@@ -510,7 +510,8 @@ train_on_inputs: false
 # Note that training loss may have an oscillating pattern with this enabled.
 group_by_length: false

-# Whether to use gradient checkpointing https://huggingface.co/docs/transformers/v4.18.0/en/performance#gradient-checkpointing
+# Whether to use gradient checkpointing. Available options are: true, false, "offload".
+# https://huggingface.co/docs/transformers/v4.18.0/en/performance#gradient-checkpointing
 gradient_checkpointing: false
 # additional kwargs to pass to the trainer for gradient checkpointing
 # gradient_checkpointing_kwargs: