@misc{indiciae37fdbcaa50e0, title = {Enhancing LLMs for Physics Problem-Solving using Reinforcement Learning with Human-AI Feedback}, author = {Avinash Anand and Kritarth Prasad and Chhavi Kirtani and Ashwin R Nair and Mohit Gupta and Saloni Garg and Anurag Gautam and Snehal Buldeo and Rajiv Ratn Shah}, year = {2024}, url = {https://arxiv.org/abs/2412.06827}, note = {Source identifier: 2412.06827} }