@misc{indiciaee0d3b8a0af02, title = {SPEQ: Offline Stabilization Phases for Efficient Q-Learning in High Update-To-Data Ratio Reinforcement Learning}, author = {Carlo Romeo and Girolamo Macaluso and Alessandro Sestini and Andrew D. Bagdanov}, year = {2025}, url = {https://arxiv.org/abs/2501.08669}, note = {Source identifier: 2501.08669} }