fix shift short attention

Former-commit-id: 9a49cce8e6f6b222f74a07bdab40efee6a77b0f1
This commit is contained in:
hiyouga
2023-10-09 17:07:46 +08:00
parent 5c4248a29c
commit e387a50475
6 changed files with 46 additions and 52 deletions

View File

@@ -29,6 +29,7 @@ def run_dpo(
dataset = preprocess_dataset(dataset, tokenizer, data_args, training_args, stage="rm")
data_collator = DPODataCollatorWithPadding(
tokenizer=tokenizer,
pad_to_multiple_of=4,
label_pad_token_id=IGNORE_INDEX if data_args.ignore_pad_token_for_loss else tokenizer.pad_token_id
)