1
0
mirror of https://github.com/gryf/coach.git synced 2025-12-17 19:20:19 +01:00

fix bug in ddpg

This commit is contained in:
Zach Dwiel
2018-02-16 20:18:03 -05:00
parent 8248caf35e
commit 5cf10e5f52
2 changed files with 2 additions and 3 deletions

View File

@@ -54,7 +54,7 @@ class DDPGAgent(ActorCriticAgent):
actions_mean = self.actor_network.online_network.predict(current_states)
critic_online_network = self.critic_network.online_network
# TODO: convert into call to predict, current method ignores lstm middleware for example
action_gradients = self.critic_network.sess.run(critic_online_network.gradients_wrt_inputs[1],
action_gradients = self.critic_network.sess.run(critic_online_network.gradients_wrt_inputs['action'],
feed_dict=critic_online_network._feed_dict({
**current_states,
'action': actions_mean,