arXiv · 2026 · Preprint

S-GRPO: Unified Post-Training for Large Vision-Language Models

Yu-Ming Yan, Kai Tang, Sihong Chen, Ke Xu, Dan Hu, Qun Yu, Pengfei Hu

Date
2026-04-17
Citations
2
arXiv
2604.16557
Cite
@misc{yan2026sgrpo,
  title = {S-GRPO: Unified Post-Training for Large Vision-Language Models},
  author = {Yu-Ming Yan and Kai Tang and Sihong Chen and Ke Xu and Dan Hu and Qun Yu and Pengfei Hu},
  year = {2026},
  eprint = {2604.16557},
  archivePrefix = {arXiv},
  url = {https://arxiv.org/abs/2604.16557},
}