@misc{indiciae6539392af76b, title = {PersoDPO: Scalable Preference Optimization for Instruction-Adherent, Persona-Grounded Dialogue via Multi-LLM Evaluation}, author = {Saleh Afzoon and MohammadHossein Ahmadi and Usman Naseem and Amin Beheshti}, year = {2026}, url = {https://arxiv.org/abs/2602.04493}, note = {Source identifier: 2602.04493} }