support DPO training (2305.18290)

Former-commit-id: 3ec4351cfd
This commit is contained in:
hiyouga
2023-08-11 03:02:53 +08:00
parent 6c32b6922b
commit abdfa26d06
34 changed files with 513 additions and 212 deletions

View File

@@ -4,7 +4,7 @@ import threading
import time
import transformers
from transformers.trainer import TRAINING_ARGS_NAME
from typing import Generator, List, Optional, Tuple
from typing import Generator, List, Tuple
from llmtuner.extras.callbacks import LogCallback
from llmtuner.extras.constants import DEFAULT_MODULE