@misc{indiciaebdab6a2aabb1, title = {Evaluating the Performance of Large Language Models via Debates}, author = {Behrad Moniri and Hamed Hassani and Edgar Dobriban}, year = {2025}, url = {https://arxiv.org/abs/2406.11044}, note = {Source identifier: 2406.11044} }