@misc{indiciaec87f81dcbf58, title = {Provably Optimal Reinforcement Learning under Safety Filtering}, author = {Donggeon David Oh and Duy P. Nguyen and Haimin Hu and Jaime Fernández Fisac}, year = {2026}, doi = {10.1609/iaseai.v2i1.43046}, url = {https://arxiv.org/abs/2510.18082}, note = {Source identifier: 2510.18082} }