[model] support yarn (#6693)

Former-commit-id: 8c412abc44a4c61b683465e36c6288580d980250
2025-12-18 12:50:38 +08:00 · 2025-01-18 13:56:09 +08:00
parent be3525910d
commit 87db2a849a
11 changed files with 84 additions and 64 deletions
--- a/src/llamafactory/webui/components/export.py
+++ b/src/llamafactory/webui/components/export.py
@@ -18,7 +18,7 @@ from ...extras.constants import PEFT_METHODS
 from ...extras.misc import torch_gc
 from ...extras.packages import is_gradio_available
 from ...train.tuner import export_model
-from ..common import GPTQ_BITS, get_save_dir
+from ..common import get_save_dir, load_config
 from ..locales import ALERTS


@@ -32,6 +32,9 @@ if TYPE_CHECKING:
    from ..engine import Engine


+GPTQ_BITS = ["8", "4", "3", "2"]
+
+
 def can_quantize(checkpoint_path: Union[str, List[str]]) -> "gr.Dropdown":
    if isinstance(checkpoint_path, list) and len(checkpoint_path) != 0:
        return gr.Dropdown(value="none", interactive=False)
@@ -54,6 +57,7 @@ def save_model(
    export_dir: str,
    export_hub_model_id: str,
 ) -> Generator[str, None, None]:
+    user_config = load_config()
    error = ""
    if not model_name:
        error = ALERTS["err_no_model"][lang]
@@ -75,6 +79,7 @@ def save_model(

    args = dict(
        model_name_or_path=model_path,
+        cache_dir=user_config.get("cache_dir", None),
        finetuning_type=finetuning_type,
        template=template,
        export_dir=export_dir,