@misc{indiciae5f0c01d4f0a9, title = {Illusions of the Gold Standard: A Large-scale Analysis of Human Evaluation Protocols for Long-form Text Generation}, author = {Katelyn Xiaoying Mei and Yi-Li Hsu and Minjoon Choi and Zongwan Cao and Chenjun Xu and Bingbing Wen and Su Lin Blodgett and Lucy Lu Wang}, year = {2026}, url = {https://arxiv.org/abs/2606.07936}, note = {Source identifier: 2606.07936} }