@misc{indiciae2dc0597f2975, title = {Reinforcement Learning for Dividend Optimization in Partially Observed Regime-Switching Diffusion Model}, author = {Zhongqin Gao and Yan Lv and Jingmin He}, year = {2026}, url = {https://arxiv.org/abs/2601.20387}, note = {Source identifier: 2601.20387} }