@misc{indiciae8fe059734e6d, title = {Aligning with Human Judgement: The Role of Pairwise Preference in Large Language Model Evaluators}, author = {Yinhong Liu and Han Zhou and Zhijiang Guo and Ehsan Shareghi and Ivan Vulić and Anna Korhonen and Nigel Collier}, year = {2025}, url = {https://arxiv.org/abs/2403.16950}, note = {Source identifier: 2403.16950} }