@misc{indiciaef76a5b870aeb, title = {Do You Really Need to Pretrain Q-Functions for Online RL Fine-Tuning?}, author = {Perry Dong and Ron Polonsky and Dorsa Sadigh and Chelsea Finn}, year = {2026}, url = {https://arxiv.org/abs/2607.27203}, note = {Source identifier: 2607.27203} }