@misc{indiciaeb9ecee29777e, title = {Q-Probe: A Lightweight Approach to Reward Maximization for Language Models}, author = {Kenneth Li and Samy Jelassi and Hugh Zhang and Sham Kakade and Martin Wattenberg and David Brandfonbrener}, year = {2024}, url = {https://arxiv.org/abs/2402.14688}, note = {Source identifier: 2402.14688} }