Fix exec, eval (#743)
* Update gemma2.py * Update llama.py * Update llama.py * Update gemma2.py * init * Update gemma2.py * Update gemma2.py * Update _utils.py * Update _utils.py * Update _utils.py * Update _utils.py * Update _utils.py * Update gemma2.py * Update gemma2.py * Update gemma2.py * All RoPE Scaling support * cleanup * Update llama.py * Update llama.py * Update _utils.py * Update _utils.py * exec * exec * Attention_Module * attention_module * imports * exec * Update llama.py * Update llama.py * boolean mask * revert masking * Update llama.py * Update save.py * Update llama.py * Update gemma2.py * Update gemma2.py * Update gemma2.py * Update utils.py * retry * Update gemma2.py * Update gemma2.py * Update gemma2.py * Update _utils.py * Update _utils.py * Update gemma2.py * Update chat_templates.py * Gemma 2 Ollama support * Update llama.py * Update llama.py * error handling * Update _utils.py * Update _utils.py * Stats for debugging * Update _utils.py * Update _utils.py * Debugging * Update tokenizer_utils.py * Update _utils.py * Update cross_entropy_loss.py * Update cross_entropy_loss.py * Update cross_entropy_loss.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Check exec, eval
This commit is contained in:
parent
c6c6eb132f
commit
b1002ba236
5 changed files with 21 additions and 9 deletions
|
|
@ -618,7 +618,11 @@ def patch_linear_scaling(
|
|||
f"from {model_filepath} import logger, "\
|
||||
f"{model_name.title()}Attention, {model_name.title()}Config"
|
||||
|
||||
function = inspect.getsource(attention_module.__init__)
|
||||
try:
|
||||
function = inspect.getsource(attention_module.__init__)
|
||||
except:
|
||||
# Most likely already patched!
|
||||
return None, None
|
||||
where = function.find("def")
|
||||
function = function.split("\n")
|
||||
function = "\n".join(x[where:] for x in function)
|
||||
|
|
|
|||
|
|
@ -281,8 +281,10 @@ class FastGemmaModel(FastLlamaModel):
|
|||
scaled_rope_module = GemmaFixedLinearScalingRotaryEmbedding,
|
||||
attention_module = GemmaAttention,
|
||||
)
|
||||
exec(function, globals())
|
||||
GemmaAttention.__init__ = eval(init_name)
|
||||
if init_name is not None:
|
||||
exec(function, globals())
|
||||
GemmaAttention.__init__ = eval(init_name)
|
||||
pass
|
||||
GemmaAttention .forward = LlamaAttention_fast_forward
|
||||
GemmaSdpaAttention .forward = LlamaAttention_fast_forward
|
||||
GemmaFlashAttention2.forward = LlamaAttention_fast_forward
|
||||
|
|
|
|||
|
|
@ -436,8 +436,10 @@ class FastGemma2Model(FastLlamaModel):
|
|||
scaled_rope_module = GemmaFixedLinearScalingRotaryEmbedding,
|
||||
attention_module = Gemma2Attention,
|
||||
)
|
||||
exec(function, globals())
|
||||
Gemma2Attention.__init__ = eval(init_name)
|
||||
if init_name is not None:
|
||||
exec(function, globals())
|
||||
Gemma2Attention.__init__ = eval(init_name)
|
||||
pass
|
||||
Gemma2Attention .forward = Gemma2Attention_fast_forward
|
||||
Gemma2SdpaAttention .forward = Gemma2Attention_fast_forward
|
||||
Gemma2FlashAttention2.forward = Gemma2Attention_fast_forward
|
||||
|
|
|
|||
|
|
@ -277,8 +277,10 @@ class FastMistralModel(FastLlamaModel):
|
|||
scaled_rope_module = LlamaLinearScalingRotaryEmbedding,
|
||||
attention_module = MistralAttention,
|
||||
)
|
||||
exec(function, globals())
|
||||
MistralAttention.__init__ = eval(init_name)
|
||||
if init_name is not None:
|
||||
exec(function, globals())
|
||||
MistralAttention.__init__ = eval(init_name)
|
||||
pass
|
||||
MistralAttention .forward = MistralAttention_fast_forward
|
||||
MistralSdpaAttention .forward = MistralAttention_fast_forward
|
||||
MistralFlashAttention2.forward = MistralAttention_fast_forward
|
||||
|
|
|
|||
|
|
@ -45,8 +45,10 @@ class FastQwen2Model(FastLlamaModel):
|
|||
scaled_rope_module = LlamaLinearScalingRotaryEmbedding,
|
||||
attention_module = Qwen2Attention,
|
||||
)
|
||||
exec(function, globals())
|
||||
Qwen2Attention.__init__ = eval(init_name)
|
||||
if init_name is not None:
|
||||
exec(function, globals())
|
||||
Qwen2Attention.__init__ = eval(init_name)
|
||||
pass
|
||||
Qwen2Attention .forward = LlamaAttention_fast_forward
|
||||
Qwen2SdpaAttention .forward = LlamaAttention_fast_forward
|
||||
Qwen2FlashAttention2.forward = LlamaAttention_fast_forward
|
||||
|
|
|
|||
Loading…
Add table
Add a link
Reference in a new issue