@misc{indiciaec9eb4502c77c, title = {Learning to Trust Bellman Updates: Selective State-Adaptive Regularization for Offline RL}, author = {Qin-Wen Luo and Ming-Kun Xie and Ye-Wen Wang and Sheng-Jun Huang}, year = {2025}, url = {https://arxiv.org/abs/2505.19923}, note = {Source identifier: 2505.19923} }