@misc{indiciae1438c5225d29, title = {s2n-bignum-bench: A practical benchmark for evaluating low-level code reasoning of LLMs}, author = {Balaji Rao and John Harrison and Soonho Kong and Juneyoung Lee and Carlo Lipizzi}, year = {2026}, url = {https://arxiv.org/abs/2603.14628}, note = {Source identifier: 2603.14628} }