* Update _utils.py * Update _utils.py * Update _utils.py * Update _utils.py * Update _utils.py * Update tokenizer_utils.py * Update tokenizer_utils.py * Update tokenizer_utils.py * update token retrieval logic (#952) * Fix DPO (#947) * Update _utils.py * Update _utils.py * Update _utils.py * Update _utils.py * Update _utils.py * Update tokenizer_utils.py * Update tokenizer_utils.py * Update tokenizer_utils.py * update hf token retrieval logic --------- Co-authored-by: Daniel Han <danielhanchen@gmail.com> * Update llama.py * get_token * Update README.md * Update gemma2.py * Update rms_layernorm.py * synchronize * Update gemma2.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * layernorm * Update rms_layernorm.py * Update gemma2.py * Update rms_layernorm.py * Update rms_layernorm.py * revert * Gemma * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update rms_layernorm.py * Update gemma2.py * Change UnslothTrainingArguments base class to SFTConfig (#979) * Cohere * Update trainer.py * Cohere * Cohere * New models * Update llama.py * Update llama.py * Update cohere.py * Update llama.py * Update cohere.py * retry * Update fast_lora.py * Update llama.py * Update fast_lora.py * Update llama.py * Update llama.py * Update cross_entropy_loss.py * _apply_lora_mlp * Update _utils.py * Gemma fixes * Update llama.py * Update flex_attention.py --------- Co-authored-by: Hafedh <70411813+not-lain@users.noreply.github.com> Co-authored-by: Tuan Pham <82665400+vTuanpham@users.noreply.github.com>
52 lines
1.7 KiB
Python
52 lines
1.7 KiB
Python
# Copyright 2023-present Daniel Han-Chen & the Unsloth team. All rights reserved.
|
|
#
|
|
# Licensed under the Apache License, Version 2.0 (the "License");
|
|
# you may not use this file except in compliance with the License.
|
|
# You may obtain a copy of the License at
|
|
#
|
|
# http://www.apache.org/licenses/LICENSE-2.0
|
|
#
|
|
# Unless required by applicable law or agreed to in writing, software
|
|
# distributed under the License is distributed on an "AS IS" BASIS,
|
|
# WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
|
|
# See the License for the specific language governing permissions and
|
|
# limitations under the License.
|
|
|
|
from .cross_entropy_loss import fast_cross_entropy_loss
|
|
from .rms_layernorm import fast_rms_layernorm
|
|
from .rope_embedding import fast_rope_embedding, inplace_rope_embedding
|
|
from .swiglu import swiglu_fg_kernel, swiglu_DWf_DW_dfg_kernel
|
|
from .geglu import (
|
|
geglu_exact_forward_kernel,
|
|
geglu_exact_backward_kernel,
|
|
geglu_approx_forward_kernel,
|
|
geglu_approx_backward_kernel,
|
|
)
|
|
from .fast_lora import (
|
|
get_lora_parameters,
|
|
get_lora_parameters_bias,
|
|
apply_lora_mlp_swiglu,
|
|
apply_lora_mlp_geglu_exact,
|
|
apply_lora_mlp_geglu_approx,
|
|
apply_lora_qkv,
|
|
apply_lora_o,
|
|
)
|
|
from .utils import fast_dequantize, fast_gemv, QUANT_STATE, fast_linear_forward, matmul_lora
|
|
|
|
from .flex_attention import (
|
|
HAS_FLEX_ATTENTION,
|
|
slow_attention_softcapping,
|
|
slow_inference_attention_softcapping,
|
|
)
|
|
|
|
if HAS_FLEX_ATTENTION:
|
|
from .flex_attention import (
|
|
FLEX_ATTENTION_PADDING,
|
|
)
|
|
pass
|
|
|
|
try:
|
|
print("🦥 Unsloth: Will patch your computer to enable 2x faster free finetuning.")
|
|
except:
|
|
print("Unsloth: Will patch your computer to enable 2x faster free finetuning.")
|
|
pass
|