@misc{indiciaece4da7068629, title = {QPLEX Decision Processes: Formulation via Nonlinear Markov Chains and Optimization via Policy Gradients}, author = {Antonius B. Dieker and Steven T. Hackman and Zitong Wang and Yunhao Yan}, year = {2026}, url = {https://arxiv.org/abs/2605.17149}, note = {Source identifier: 2605.17149} }