@misc{indiciae0b0ba60636f3, title = {Q-Flow: Stable and Expressive Reinforcement Learning with Flow-Based Policy}, author = {JaeHyeok Doo and Byeongguk Jeon and Seonghyeon Ye and Kimin Lee and Minjoon Seo}, year = {2026}, url = {https://arxiv.org/abs/2605.13435}, note = {Source identifier: 2605.13435} }