@misc{indiciae6fca915e64ce, title = {ScholarGym: Benchmarking Large Language Model Capabilities in the Information-Gathering Stage of Deep Research}, author = {Hao Shen and Hang Yang and Zhouhong Gu and Weili Han}, year = {2026}, url = {https://arxiv.org/abs/2601.21654}, note = {Source identifier: 2601.21654} }