Update tokenizer_utils.py

This commit is contained in:
Daniel Han 2024-09-23 00:04:33 -07:00
commit 1b8ef43c14

View file

@ -942,7 +942,7 @@ def fix_untrained_tokens(model, tokenizer, train_dataset, eps = 1e-16):
pass
raise ValueError(
f'Unsloth: Untrained tokens of [{list(final_bad_items)}] found, but embed_tokens & lm_head not trainable, causing NaNs. '\
f'Unsloth: Untrained tokens of [{list(set(final_bad_items))}] found, but embed_tokens & lm_head not trainable, causing NaNs. '\
'Restart then add `embed_tokens` & `lm_head` to '\
'`FastLanguageModel.get_peft_model(target_modules = [..., "embed_tokens", "lm_head",]). `'\
'Are you using the `base` model? Instead, use the `instruct` version to silence this warning.',