@misc{indiciae075db553706b, title = {Are Large Language Models Reliable Reviewers? A Benchmark for Error Detection in Financial Documents}, author = {Ying He and Zhouhong Gu and Zhecheng Hu and Yubo Zhou and Hao Shen and Jiaqing Liang and Zhaoqian Dai and Shuguang Ma and Fei Yu and Yanghua Xiao and Zhixu Li}, year = {2026}, url = {https://arxiv.org/abs/2608.12342}, note = {Source identifier: 2608.12342} }