@article{Liu2025, 
author = {Lei Liu and Mengmeng Hao and Jinde Cao},
title = {Dynamic hedging of 50ETF options using Proximal Policy Optimization},
year = {2025},
journal = {Journal of Automation and Intelligence},
volume = {4},
number = {3},
pages = {198-206},
keywords = {B–S model, Option hedging, Reinforcement learning, 50ETF, Proximal Policy Optimization (PPO)},
url = {https://www.sciopen.com/article/10.1016/j.jai.2025.04.001},
doi = {10.1016/j.jai.2025.04.001},
abstract = {This paper employs the PPO (Proximal Policy Optimization) algorithm to study the risk hedging problem of the Shanghai Stock Exchange (SSE) 50ETF options. First, the action and state spaces were designed based on the characteristics of the hedging task, and a reward function was developed according to the cost function of the options. Second, combining the concept of curriculum learning, the agent was guided to adopt a simulated-to-real learning approach for dynamic hedging tasks, reducing the learning difficulty and addressing the issue of insufficient option data. A dynamic hedging strategy for 50ETF options was constructed. Finally, numerical experiments demonstrate the superiority of the designed algorithm over traditional hedging strategies in terms of hedging effectiveness.}
}