don't use mask expansion for inference (#392)

2023-08-14 20:52:54 -04:00
parent 41ecb451c2
commit 1687be6a35
3 changed files with 6 additions and 2 deletions
--- a/examples/llama-2/lora.yml
+++ b/examples/llama-2/lora.yml
@@ -2,6 +2,7 @@ base_model: meta-llama/Llama-2-7b-hf
 base_model_config: meta-llama/Llama-2-7b-hf
 model_type: LlamaForCausalLM
 tokenizer_type: LlamaTokenizer
+is_llama_derived_model: true

 load_in_8bit: true
 load_in_4bit: false
--- a/examples/llama-2/qlora.yml
+++ b/examples/llama-2/qlora.yml
@@ -2,6 +2,7 @@ base_model: meta-llama/Llama-2-7b-hf
 base_model_config: meta-llama/Llama-2-7b-hf
 model_type: LlamaForCausalLM
 tokenizer_type: LlamaTokenizer
+is_llama_derived_model: true

 load_in_8bit: false
 load_in_4bit: true