CPPO: Continual Learning for Reinforcement Learning with Human Feedback
- International Conference on Learning Representations
- 53
@article{zhang2024cppo,
title = {CPPO: Continual Learning for Reinforcement Learning with Human Feedback},
author = {Han Zhang and Yu Lei and Lin Gui and Min Yang and Yulan He and Hui Wang and Ruifeng Xu},
year = {2024},
journal = {International Conference on Learning Representations},
url = {https://www.semanticscholar.org/paper/c2f1fbc5e829a8b667470bb4f61698f25210442b},
}