@misc{indiciae3d8ab76de8a9, title = {VPP: Virtual Pipeline Parallelism for Efficient Chunked Prefill in Long-Context LLM Inference}, author = {Yan Shi and Xiaochao Wang and Jingchun Gao and Jintao Luo and Xinyi Zhou and Feng Liu and Kui Luo and Xushi Li and Xinjie Guo and Liangjun Feng}, year = {2026}, url = {https://arxiv.org/abs/2608.26523}, note = {Source identifier: 2608.26523} }