@misc{indiciaef725e2afa416, title = {SLO-Aware Compute Resource Allocation for Prefill-Decode Disaggregated LLM Inference}, author = {Luchang Li and Dongfang Li and Bozhao Gong and Yu Zhang}, year = {2026}, url = {https://arxiv.org/abs/2603.04716}, note = {Source identifier: 2603.04716} }