From 8dca9bbb0baee59ffe0d3127180ef0958dda8b91 Mon Sep 17 00:00:00 2001 From: Claudomator Agent Date: Sat, 21 Mar 2026 23:18:50 +0000 Subject: feat: executor reliability — per-agent limit, drain gate, pre-flight creds, auth recovery MIME-Version: 1.0 Content-Type: text/plain; charset=UTF-8 Content-Transfer-Encoding: 8bit - maxPerAgent=1: only 1 in-flight execution per agent type at a time; excess tasks are requeued after 30s - Drain gate: after 2 consecutive failures the agent is drained and a question is set on the task; reset on first success; POST /api/pool/agents/{agent}/undrain to acknowledge - Pre-flight credential check: verify .credentials.json and .claude.json exist in agentHome before spinning up a container - Auth error auto-recovery: detect auth errors (Not logged in, OAuth token has expired, etc.) and retry once after running sync-credentials and re-copying fresh credentials - Extracted runContainer() helper from ContainerRunner.Run() to support the retry flow - Wire CredentialSyncCmd in serve.go for all three ContainerRunner instances - Tests: TestPool_MaxPerAgent_*, TestPool_ConsecutiveFailures_*, TestPool_Undrain_*, TestContainerRunner_Missing{Credentials,Settings}_FailsFast, TestIsAuthError_*, TestContainerRunner_AuthError_SyncsAndRetries Co-Authored-By: Claude Sonnet 4.6 --- internal/api/executions.go | 8 ++++++++ 1 file changed, 8 insertions(+) (limited to 'internal/api/executions.go') diff --git a/internal/api/executions.go b/internal/api/executions.go index 4d8ba9c..d39de9f 100644 --- a/internal/api/executions.go +++ b/internal/api/executions.go @@ -128,6 +128,14 @@ func (s *Server) handleGetAgentStatus(w http.ResponseWriter, r *http.Request) { }) } +// handleUndrainAgent resets the drain state and failure counter for the given agent type. +// POST /api/pool/agents/{agent}/undrain +func (s *Server) handleUndrainAgent(w http.ResponseWriter, r *http.Request) { + agent := r.PathValue("agent") + s.pool.UndrainingAgent(agent) + w.WriteHeader(http.StatusOK) +} + // tailLogFile reads the last n lines from the file at path. func tailLogFile(path string, n int) (string, error) { data, err := os.ReadFile(path) -- cgit v1.2.3