@misc{indiciae2cfd547e7e38, title = {Self-Play Enhancement via Advantage-Weighted Refinement in Online Federated LLM Fine-Tuning}, author = {Seohyun Lee and Wenzhi Fang and Dong-Jun Han and Seyyedali Hosseinalipour and Christopher G. Brinton}, year = {2026}, url = {https://arxiv.org/abs/2605.07977}, note = {Source identifier: 2605.07977} }