@misc{indiciae870b60628545, title = {Control Reinforcement Learning: Interpretable Token-Level Steering of LLMs via Sparse Autoencoder Features}, author = {Seonglae Cho and Zekun Wu and Adriano Koshiyama}, year = {2026}, url = {https://arxiv.org/abs/2602.10437}, note = {Source identifier: 2602.10437} }