<?xml version="1.0" encoding="UTF-8"?>
<urlset xmlns="http://www.sitemaps.org/schemas/sitemap/0.9">
<url>
<loc>https://rl4rvr.shannon.id/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/about/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/why-rl-for-robotics/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/mathematical-toolkit/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/multi-armed-bandits/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/markov-decision-processes/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/dynamic-programming/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/monte-carlo-and-td/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/traces-planning-mcts/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/function-approximation/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/deep-value-methods/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/policy-gradients/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/off-policy-continuous-control/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/model-based-rl/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/robot-as-environment/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/four-curses/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/sim-to-real/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/imitation-and-offline-rl/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/motor-skill-representations/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/learning-locomotion/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/navigation-and-mobile-manipulation/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/learning-manipulation/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/frontiers/</loc>
</url>
<url>
<loc>https://rl4rvr.shannon.id/chapters/capstone/</loc>
</url>
</urlset>
