judge_coding_task_completion
Validate coding task completion by reviewing requirements, implementation, and testing. Approve only when all criteria are met; otherwise, list required improvements.
Instructions
Judge Coding Task Completion
Description
Final validation gate before declaring a task complete. Called when workflow_guidance.next_tool == "judge_coding_task_completion".
Critical Tool Warning
Skipping this tool causes severe token inefficiency and wasted iterations.
Always invoke this tool at the appropriate stage to avoid extreme token loss and redundant processing.
Do not rely on assistant memory for identifiers. Always pass the exact
task_idand recover it viaget_current_coding_taskif missing.
Prerequisites
Plan approved via
judge_coding_plan, code approved viajudge_code_change, tests approved viajudge_testing_implementation
Args
task_id: string — Task UUID (required)completion_summary: string — Summary of implemented work (required)requirements_met: list[string] — Requirements satisfied (required)implementation_details: string — Key implementation details (required)remaining_work: list[string] — Open items if any (optional)quality_notes: string — Quality/standards notes (optional)testing_status: string — Testing status summary (optional)
Returns
Response JSON schema (TaskCompletionResult):
{
"$defs": {
"LibraryPlanItem": {
"properties": {
"purpose": {
"description": "Non-domain concern or integration point this library addresses",
"title": "Purpose",
"type": "string"
},
"selection": {
"description": "Chosen library or internal utility (name and optional version)",
"title": "Selection",
"type": "string"
},
"source": {
"description": "Source of solution: 'internal' for repo utility, 'external' for well-known library, 'custom' for in-house code",
"title": "Source",
"type": "string"
},
"justification": {
"default": "",
"description": "One-line rationale for the selection and any trade-offs",
"title": "Justification",
"type": "string"
}
},
"required": [
"purpose",
"selection",
"source"
],
"title": "LibraryPlanItem",
"type": "object"
},
"PlanRequiredField": {
"description": "Specification for a required field in judge_coding_plan.",
"properties": {
"name": {
"description": "Field name in the judge_coding_plan tool",
"title": "Name",
"type": "string"
},
"type": {
"description": "Expected data type (string, list[str], list[dict], etc.)",
"title": "Type",
"type": "string"
},
"description": {
"description": "What this field should contain",
"title": "Description",
"type": "string"
},
"required": {
"description": "Whether this field is required",
"title": "Required",
"type": "boolean"
},
"conditional_on": {
"anyOf": [
{
"type": "string"
},
{
"type": "null"
}
],
"default": null,
"description": "Task metadata field this requirement depends on (e.g., 'design_patterns_enforcement')",
"title": "Conditional On"
},
"example_value": {
"anyOf": [
{
"type": "string"
},
{
"type": "null"
}
],
"default": null,
"description": "Example of what this field should contain",
"title": "Example Value"
}
},
"required": [
"name",
"type",
"description",
"required"
],
"title": "PlanRequiredField",
"type": "object"
},
"RequirementsVersion": {
"description": "A version of user requirements with timestamp and source.",
"properties": {
"content": {
"title": "Content",
"type": "string"
},
"source": {
"title": "Source",
"type": "string"
},
"timestamp": {
"title": "Timestamp",
"type": "integer"
}
},
"required": [
"content",
"source"
],
"title": "RequirementsVersion",
"type": "object"
},
"ResearchScope": {
"description": "Research scope enum for workflow-driven research validation.\n\nDetermines the depth and requirements for research validation:\n- NONE: No research required for this task complexity\n- LIGHT: Light research required (1+ authoritative domain source)\n- DEEP: Deep research required (2+ authoritative domain sources)",
"enum": [
"none",
"light",
"deep"
],
"title": "ResearchScope",
"type": "string"
},
"ReuseComponent": {
"properties": {
"path": {
"description": "Repository path to the reusable component",
"title": "Path",
"type": "string"
},
"purpose": {
"default": "",
"description": "What part of the task this component will support",
"title": "Purpose",
"type": "string"
},
"notes": {
"default": "",
"description": "Any integration notes or caveats",
"title": "Notes",
"type": "string"
}
},
"required": [
"path"
],
"title": "ReuseComponent",
"type": "object"
},
"TaskMetadata": {
"description": "Lightweight metadata for coding tasks that flows with memory layer.\n\nThis model serves as the foundation for the enhanced workflow v3 system,\nreplacing session-based tracking with task-centric approach.",
"properties": {
"task_id": {
"description": "IMMUTABLE: Auto-generated UUID, primary key for memory storage",
"title": "Task Id",
"type": "string"
},
"created_at": {
"description": "IMMUTABLE: Task creation timestamp (epoch seconds)",
"title": "Created At",
"type": "integer"
},
"title": {
"description": "Display title for coding task (updatable)",
"title": "Title",
"type": "string"
},
"description": {
"description": "Detailed coding task description (updatable)",
"title": "Description",
"type": "string"
},
"user_requirements": {
"default": "",
"description": "Current coding requirements (updatable)",
"title": "User Requirements",
"type": "string"
},
"state": {
"$ref": "#/$defs/TaskState",
"default": "created",
"description": "Current task state (updatable, follows TaskState transitions)"
},
"task_size": {
"$ref": "#/$defs/TaskSize",
"description": "Task size classification for workflow optimization (XS=simple fixes, S=minor features, M=standard, L=complex, XL=major changes)"
},
"user_requirements_history": {
"description": "History of requirements changes",
"items": {
"$ref": "#/$defs/RequirementsVersion"
},
"title": "User Requirements History",
"type": "array"
},
"accumulated_diff": {
"additionalProperties": true,
"description": "Code changes accumulated over time",
"title": "Accumulated Diff",
"type": "object"
},
"modified_files": {
"description": "List of file paths that were created or modified during task implementation",
"items": {
"type": "string"
},
"title": "Modified Files",
"type": "array"
},
"test_files": {
"description": "List of test file paths that were created during testing phase",
"items": {
"type": "string"
},
"title": "Test Files",
"type": "array"
},
"test_status": {
"additionalProperties": {
"type": "string"
},
"description": "Status of different test types (unit, integration, e2e, etc.)",
"title": "Test Status",
"type": "object"
},
"updated_at": {
"description": "Last update timestamp (epoch seconds)",
"title": "Updated At",
"type": "integer"
},
"tags": {
"description": "Coding-related tags",
"items": {
"type": "string"
},
"title": "Tags",
"type": "array"
},
"problem_domain": {
"default": "",
"description": "Concise statement of the problem domain and scope for this task",
"title": "Problem Domain",
"type": "string"
},
"problem_non_goals": {
"description": "Explicit non-goals/boundaries to prevent scope creep and re-solving commodity concerns",
"items": {
"type": "string"
},
"title": "Problem Non Goals",
"type": "array"
},
"library_plan": {
"description": "Planned libraries/utilities per purpose; prefer internal reuse and well-known libraries; custom code only with justification",
"items": {
"$ref": "#/$defs/LibraryPlanItem"
},
"title": "Library Plan",
"type": "array"
},
"internal_reuse_components": {
"description": "Existing repository components/utilities to reuse with paths and purposes",
"items": {
"$ref": "#/$defs/ReuseComponent"
},
"title": "Internal Reuse Components",
"type": "array"
},
"research_required": {
"anyOf": [
{
"type": "boolean"
},
{
"type": "null"
}
],
"default": null,
"description": "Whether research is required by workflow guidance (None=undetermined, True=required, False=optional)",
"title": "Research Required"
},
"research_scope": {
"$ref": "#/$defs/ResearchScope",
"default": "none",
"description": "Research scope determined by workflow: none|light|deep"
},
"research_completed": {
"anyOf": [
{
"type": "integer"
},
{
"type": "null"
}
],
"default": null,
"description": "Epoch seconds when research validation passed",
"title": "Research Completed"
},
"research_rationale": {
"default": "",
"description": "Explanation of why research was required and how the scope was determined",
"title": "Research Rationale",
"type": "string"
},
"expected_url_count": {
"anyOf": [
{
"type": "integer"
},
{
"type": "null"
}
],
"default": null,
"description": "LLM-determined expected number of research URLs based on task complexity",
"title": "Expected Url Count"
},
"minimum_url_count": {
"anyOf": [
{
"type": "integer"
},
{
"type": "null"
}
],
"default": null,
"description": "LLM-determined minimum acceptable URL count for adequate research",
"title": "Minimum Url Count"
},
"url_requirement_reasoning": {
"default": "",
"description": "LLM-generated explanation of why specific URL count is needed for this task",
"title": "Url Requirement Reasoning",
"type": "string"
},
"research_complexity_analysis": {
"anyOf": [
{
"additionalProperties": true,
"type": "object"
},
{
"type": "null"
}
],
"default": null,
"description": "Detailed complexity analysis factors from LLM (domain, tech maturity, integration scope, etc.)",
"title": "Research Complexity Analysis"
},
"internal_research_required": {
"anyOf": [
{
"type": "boolean"
},
{
"type": "null"
}
],
"default": null,
"description": "Whether internal codebase research is needed (None=undetermined, True=required, False=not needed)",
"title": "Internal Research Required"
},
"related_code_snippets": {
"description": "Related code snippets from the codebase that are relevant to this task",
"items": {
"type": "string"
},
"title": "Related Code Snippets",
"type": "array"
},
"risk_assessment_required": {
"anyOf": [
{
"type": "boolean"
},
{
"type": "null"
}
],
"default": null,
"description": "Whether risk assessment is needed (None=undetermined, True=required, False=not needed)",
"title": "Risk Assessment Required"
},
"identified_risks": {
"description": "Areas that could be harmed by the proposed changes",
"items": {
"type": "string"
},
"title": "Identified Risks",
"type": "array"
},
"risk_mitigation_strategies": {
"description": "Strategies to mitigate identified risks",
"items": {
"type": "string"
},
"title": "Risk Mitigation Strategies",
"type": "array"
},
"design_patterns_enforcement": {
"anyOf": [
{
"type": "boolean"
},
{
"type": "null"
}
],
"default": null,
"description": "Whether design patterns are required for this task (None=undetermined, True=required, False=not needed)",
"title": "Design Patterns Enforcement"
},
"plan_approved_at": {
"anyOf": [
{
"type": "integer"
},
{
"type": "null"
}
],
"default": null,
"description": "Timestamp when plan was approved by judge_coding_plan (None=not approved)",
"title": "Plan Approved At"
},
"plan_rejection_count": {
"default": 0,
"description": "Number of times the plan has been rejected (max 1 allowed)",
"title": "Plan Rejection Count",
"type": "integer"
},
"code_approved_files": {
"additionalProperties": {
"type": "integer"
},
"description": "Dictionary mapping file paths to approval timestamps from judge_code_change",
"title": "Code Approved Files",
"type": "object"
},
"testing_approved_at": {
"anyOf": [
{
"type": "integer"
},
{
"type": "null"
}
],
"default": null,
"description": "Timestamp when testing was approved by judge_testing_implementation (None=not approved)",
"title": "Testing Approved At"
},
"all_approvals_validated": {
"default": false,
"description": "Whether all required approvals (plan, code, testing) have been validated",
"title": "All Approvals Validated",
"type": "boolean"
}
},
"required": [
"title",
"description",
"task_size"
],
"title": "TaskMetadata",
"type": "object"
},
"TaskSize": {
"description": "Task size classification for workflow optimization.\n\nSizes are based on estimated complexity and time requirements:\n- XS: Extra Small - Simple fixes, typos, minor config changes (< 30 minutes)\n- S: Small - Minor features, simple refactoring (30 minutes - 2 hours)\n- M: Medium - Standard features, moderate complexity (2-8 hours) - DEFAULT\n- L: Large - Complex features, multiple components (1-3 days)\n- XL: Extra Large - Major system changes, architectural updates (3+ days)\n\nThis classification determines planning complexity and validation depth:\n- XS/S: Basic planning requirements, streamlined validation\n- M: Standard planning and validation\n- L/XL: Comprehensive planning with enhanced validation (library plans, risk assessment, design patterns)\n\nAll tasks follow the unified workflow: CREATED \u2192 PLANNING \u2192 PLAN_APPROVED \u2192 IMPLEMENTING \u2192 REVIEW_READY \u2192 TESTING \u2192 COMPLETED",
"enum": [
"xs",
"s",
"m",
"l",
"xl"
],
"title": "TaskSize",
"type": "string"
},
"TaskState": {
"description": "Coding task state enum with well-documented transitions.\n\nState Transitions:\n- CREATED \u2192 PLANNING: Task created, ready for planning phase (XS/S may skip to IMPLEMENTING)\n- PLANNING \u2192 PLAN_PENDING_APPROVAL: Plan created, awaiting user approval\n- PLAN_PENDING_APPROVAL \u2192 PLANNING: User requests plan changes\n- PLAN_PENDING_APPROVAL \u2192 PLAN_APPROVED: User approves plan\n- PLAN_APPROVED \u2192 IMPLEMENTING: Implementation phase started\n- IMPLEMENTING \u2192 IMPLEMENTING: Multiple code changes during implementation\n- IMPLEMENTING \u2192 REVIEW_READY: Implementation complete, ready for code review\n- REVIEW_READY \u2192 TESTING: Code review approved, ready for testing validation\n- TESTING \u2192 TESTING: Multiple test iterations\n- TESTING \u2192 COMPLETED: All tests validated; task completed successfully\n- Any state \u2192 BLOCKED: Task blocked by external dependencies\n- Any state \u2192 CANCELLED: Task cancelled\n- BLOCKED \u2192 Previous state: Unblocked, return to previous state\n\nUsage:\n- CREATED: Default state for new tasks, all tasks proceed to planning (unified workflow)\n- PLANNING: Planning phase in progress (set when planning starts)\n- PLAN_PENDING_APPROVAL: Plan created, awaiting user approval and potential iteration\n- PLAN_APPROVED: Plan validated and approved (set by judge_coding_plan)\n- IMPLEMENTING: Implementation phase in progress (set when coding starts)\n- REVIEW_READY: Implementation complete and ready for code review\n- TESTING: Testing/validation phase after code review approval\n- COMPLETED: Task completed successfully (set by judge_coding_task_completion)\n- BLOCKED: Task blocked by external dependencies (manual override)\n- CANCELLED: Task cancelled (manual override)",
"enum": [
"created",
"planning",
"plan_pending_approval",
"plan_approved",
"implementing",
"testing",
"review_ready",
"completed",
"blocked",
"cancelled"
],
"title": "TaskState",
"type": "string"
},
"WorkflowGuidance": {
"description": "Canonical workflow guidance model used across the system.\n\nReturned by tools to provide consistent next steps and instructions for\nthe coding assistant. This is the single source of truth for the\nWorkflowGuidance schema.",
"properties": {
"next_tool": {
"anyOf": [
{
"type": "string"
},
{
"type": "null"
}
],
"default": null,
"description": "Next tool to call, or None if workflow complete",
"title": "Next Tool"
},
"reasoning": {
"default": "",
"description": "Clear explanation of why this tool should be used next",
"title": "Reasoning",
"type": "string"
},
"preparation_needed": {
"description": "List of things that need to be prepared before calling the recommended tool",
"items": {
"type": "string"
},
"title": "Preparation Needed",
"type": "array"
},
"guidance": {
"default": "",
"description": "Detailed step-by-step guidance for the AI assistant",
"title": "Guidance",
"type": "string"
},
"research_required": {
"anyOf": [
{
"type": "boolean"
},
{
"type": "null"
}
],
"default": null,
"description": "Whether research is required for this task (only determined for new CREATED tasks)",
"title": "Research Required"
},
"research_scope": {
"anyOf": [
{
"type": "string"
},
{
"type": "null"
}
],
"default": null,
"description": "Research scope: 'none', 'light', or 'deep' (only determined for new CREATED tasks)",
"title": "Research Scope"
},
"research_rationale": {
"anyOf": [
{
"type": "string"
},
{
"type": "null"
}
],
"default": null,
"description": "Explanation of research requirements (only determined for new CREATED tasks)",
"title": "Research Rationale"
},
"internal_research_required": {
"anyOf": [
{
"type": "boolean"
},
{
"type": "null"
}
],
"default": null,
"description": "Whether internal codebase analysis is needed (only determined for new CREATED tasks)",
"title": "Internal Research Required"
},
"risk_assessment_required": {
"anyOf": [
{
"type": "boolean"
},
{
"type": "null"
}
],
"default": null,
"description": "Whether risk assessment is needed (only determined for new CREATED tasks)",
"title": "Risk Assessment Required"
},
"design_patterns_enforcement": {
"anyOf": [
{
"type": "boolean"
},
{
"type": "null"
}
],
"default": null,
"description": "Whether design patterns are required (only determined for new CREATED tasks)",
"title": "Design Patterns Enforcement"
},
"plan_required_fields": {
"description": "Structured specification of required fields for judge_coding_plan tool",
"items": {
"$ref": "#/$defs/PlanRequiredField"
},
"title": "Plan Required Fields",
"type": "array"
}
},
"title": "WorkflowGuidance",
"type": "object"
}
},
"properties": {
"approved": {
"description": "Whether the task completion is approved",
"title": "Approved",
"type": "boolean"
},
"feedback": {
"description": "Detailed feedback about the completion validation",
"title": "Feedback",
"type": "string"
},
"required_improvements": {
"description": "List of required improvements if not approved",
"items": {
"type": "string"
},
"title": "Required Improvements",
"type": "array"
},
"current_task_metadata": {
"$ref": "#/$defs/TaskMetadata",
"description": "ALWAYS current state of task metadata after operation"
},
"workflow_guidance": {
"$ref": "#/$defs/WorkflowGuidance",
"description": "LLM-generated next steps and instructions (or workflow complete)"
}
},
"required": [
"approved",
"feedback",
"current_task_metadata",
"workflow_guidance"
],
"title": "TaskCompletionResult",
"type": "object"
}Notes
The AI coding assistant MUST NOT present or claim task completion, or provide a final completion summary to the user, without successfully calling this tool and receiving approval.
Always use the exact
task_id; if missing due to memory limits, recover it viaget_current_coding_task.
Input Schema
| Name | Required | Description | Default |
|---|---|---|---|
| task_id | Yes | ||
| quality_notes | No | ||
| remaining_work | No | ||
| testing_status | No | ||
| requirements_met | Yes | ||
| completion_summary | Yes | ||
| implementation_details | Yes |
Output Schema
| Name | Required | Description | Default |
|---|---|---|---|
| approved | Yes | Whether the task completion is approved | |
| feedback | Yes | Detailed feedback about the completion validation | |
| workflow_guidance | Yes | LLM-generated next steps and instructions (or workflow complete) | |
| current_task_metadata | Yes | ALWAYS current state of task metadata after operation | |
| required_improvements | No | List of required improvements if not approved |