@misc{indiciae2e76b6b209b6, title = {APLOT: Robust Reward Modeling via Adaptive Preference Learning with Optimal Transport}, author = {Zhuo Li and Yuege Feng and Dandan Guo and Jinpeng Hu and Anningzhe Gao and Xiang Wan}, year = {2025}, url = {https://arxiv.org/abs/2510.10963}, note = {Source identifier: 2510.10963} }