@misc{indiciaefd5ecd02a533, title = {When Can We Trust LLMs in Mental Health? Large-Scale Benchmarks for Reliable LLM Evaluation}, author = {Abeer Badawi and Elahe Rahimi and Md Tahmid Rahman Laskar and Sheri Grach and Lindsay Bertrand and Lames Danok and Jimmy Huang and Frank Rudzicz and Elham Dolatabadi}, year = {2025}, url = {https://arxiv.org/abs/2510.19032}, note = {Source identifier: 2510.19032} }