@misc{indiciaec25e0c003885, title = {Policy Gradient Method for LQG Control via Input-Output-History Representation: Convergence to \$O(ε)\$-Stationary Points}, author = {Tomonori Sadamoto and Takashi Tanaka}, year = {2025}, url = {https://arxiv.org/abs/2510.19141}, note = {Source identifier: 2510.19141} }