# Conflicts: # apps/desktop/e2e/archived-hidden-session-recoverable.spec.ts # apps/desktop/e2e/bot-chat-message-agent-friendly-name.spec.ts # apps/desktop/e2e/bot-mailbox-unreadable-ticket.spec.ts # apps/desktop/e2e/bot-mode-roster-localized.spec.ts # apps/desktop/e2e/bot-mode-row-click-mirrors-registry.spec.ts # apps/desktop/e2e/bot-mode-tab-shows-bot-name.spec.ts # apps/desktop/e2e/bot-roster-group-row-organisation.spec.ts # apps/desktop/e2e/bot-roster-ignores-infra-dirs.spec.ts # apps/desktop/e2e/bot-roster-timestamp-meta.spec.ts # apps/desktop/e2e/bot-roster-user-sections.spec.ts # apps/desktop/e2e/bot-routines-pane-narrow.spec.ts # apps/desktop/e2e/bot-row-open-recent-session.spec.ts # apps/desktop/e2e/bot-tile-ignores-ambient-composer-model.spec.ts # apps/desktop/e2e/group-composer-auto-grow.spec.ts # apps/desktop/e2e/group-create-gate-remote-roster.spec.ts # apps/desktop/e2e/group-prompt-renamed-primary-handle.spec.ts # apps/desktop/e2e/hosted-room-backend-continuity.spec.ts # apps/desktop/e2e/hosted-room-legacy-store-migration.spec.ts # apps/desktop/e2e/settings-scope-chips-bot-title.spec.ts # apps/desktop/e2e/worktree-branch-status.spec.ts # apps/desktop/electron/backend-probes.test.ts # apps/desktop/electron/connection-apply.test.ts # apps/desktop/electron/desktop-electron-pin.test.ts # apps/desktop/electron/desktop-uninstall.test.ts # apps/desktop/electron/gateway-file-download-transport.test.ts # apps/desktop/electron/gateway-stop-before-update.test.ts # apps/desktop/electron/github-api-auth.test.ts # apps/desktop/electron/registry-primary-profile-scope.test.ts # apps/desktop/electron/update-api-check.test.ts # apps/desktop/electron/update-handoff-marker.test.ts # apps/desktop/electron/venv-blocker-scan.test.ts # apps/desktop/scripts/after-extract.test.mjs # apps/desktop/scripts/local-pack-publish.test.mjs # apps/desktop/scripts/tasks-scroll.test.mjs # apps/desktop/src/app/settings/model-settings.test.tsx # apps/desktop/src/app/updates-overlay.blockers.test.tsx # apps/desktop/src/components/desktop-install-overlay.test.tsx # apps/desktop/src/lib/update-copy.test.ts # scripts/ci/check_os_marker_fakes.py # tests-js/desktop-mac-usage-descriptions.test.ts # tests-js/node-engine-alignment.test.ts # tests/agent/lsp/test_install_and_lint_fixes.py # tests/agent/test_command_token_source.py # tests/agent/test_compression_boundary_hook.py # tests/agent/test_create_openai_client_ssl_verify.py # tests/agent/test_custom_provider_ca_probes.py # tests/agent/test_endpoint_blackhole.py # tests/agent/test_estimator_parity.py # tests/agent/test_in_place_compaction.py # tests/agent/test_moa_loop_mode.py # tests/agent/test_model_metadata.py # tests/agent/test_skill_session_platform_gate.py # tests/agent/test_skill_utils.py # tests/agent/test_ssl_ca_guard.py # tests/computer_use/test_doctor.py # tests/cron/test_codex_execution_paths.py # tests/cron/test_cron_bot_chat_delivery.py # tests/cron/test_cron_script.py # tests/cron/test_media_delivery_parity.py # tests/cron/test_misfire_catchup.py # tests/cron/test_parallel_pool.py # tests/cron/test_recurring_eagain_redispatch.py # tests/gateway/test_choice_picker.py # tests/gateway/test_control_socket_windows_live.py # tests/gateway/test_dingtalk.py # tests/gateway/test_feishu.py # tests/gateway/test_feishu_onboard.py # tests/gateway/test_gateway_shutdown.py # tests/gateway/test_matrix.py # tests/gateway/test_model_command_custom_providers.py # tests/gateway/test_reasoning_command.py # tests/gateway/test_runtime_footer.py # tests/gateway/test_session.py # tests/gateway/test_session_hygiene.py # tests/gateway/test_status.py # tests/gateway/test_teams.py # tests/gateway/test_turn_lease.py # tests/gateway/test_whatsapp_connect.py # tests/hermes_cli/test_approvals_command.py # tests/hermes_cli/test_auth_store_lock_concurrent.py # tests/hermes_cli/test_backup.py # tests/hermes_cli/test_banner_git_state.py # tests/hermes_cli/test_certifi_repair.py # tests/hermes_cli/test_cmd_update.py # tests/hermes_cli/test_compat_manifest_targets.py # tests/hermes_cli/test_computer_use_cli.py # tests/hermes_cli/test_cpr_local_leak.py # tests/hermes_cli/test_dashboard_auth_gate.py # tests/hermes_cli/test_dashboard_procs_kill_grace.py # tests/hermes_cli/test_desktop_lifecycle_windows_live.py # tests/hermes_cli/test_doctor.py # tests/hermes_cli/test_doctor_command_install.py # tests/hermes_cli/test_fleet_config_migration_windows_live.py # tests/hermes_cli/test_gateway.py # tests/hermes_cli/test_gateway_platform_gating.py # tests/hermes_cli/test_gateway_restart_loop.py # tests/hermes_cli/test_gateway_task_probe.py # tests/hermes_cli/test_gateway_wsl.py # tests/hermes_cli/test_gui_command.py # tests/hermes_cli/test_install_cua_driver.py # tests/hermes_cli/test_kanban_db.py # tests/hermes_cli/test_lazy_command_exports.py # tests/hermes_cli/test_lazy_refresh_venv_repair.py # tests/hermes_cli/test_linux_desktop_entry.py # tests/hermes_cli/test_local_runtime.py # tests/hermes_cli/test_local_runtime_updates.py # tests/hermes_cli/test_managed_uv.py # tests/hermes_cli/test_mcp_reload_confirm_gate.py # tests/hermes_cli/test_nous_subscription.py # tests/hermes_cli/test_npm_engine.py # tests/hermes_cli/test_personality_none.py # tests/hermes_cli/test_pet_toggle.py # tests/hermes_cli/test_plan_reconciliation_windows_live.py # tests/hermes_cli/test_plugin_event_bus.py # tests/hermes_cli/test_plugin_manifest_v2.py # tests/hermes_cli/test_plugin_packs.py # tests/hermes_cli/test_plugins_cmd.py # tests/hermes_cli/test_plugins_cmd_enable_disable_nested.py # tests/hermes_cli/test_process_identity.py # tests/hermes_cli/test_profiles.py # tests/hermes_cli/test_profiles_sidebar_cache.py # tests/hermes_cli/test_pty_bridge.py # tests/hermes_cli/test_resolve_turn_limit.py # tests/hermes_cli/test_serve_runtime_inventory.py # tests/hermes_cli/test_session_vacuum_config.py # tests/hermes_cli/test_set_config_value.py # tests/hermes_cli/test_signal_handler_kanban_worker.py # tests/hermes_cli/test_slash_confirm_windows.py # tests/hermes_cli/test_stale_pid_guard.py # tests/hermes_cli/test_startup_fast_guards.py # tests/hermes_cli/test_status.py # tests/hermes_cli/test_telegram_managed_bot.py # tests/hermes_cli/test_tools_config.py # tests/hermes_cli/test_update_apply_shallow_count.py # tests/hermes_cli/test_update_autostash.py # tests/hermes_cli/test_update_concurrent_quarantine.py # tests/hermes_cli/test_update_fetch_failure_classifier.py # tests/hermes_cli/test_update_fleet_probe_resume_token.py # tests/hermes_cli/test_update_handoff_backend_reap.py # tests/hermes_cli/test_update_handoff_desktop_rebuild.py # tests/hermes_cli/test_update_head_moved_gate.py # tests/hermes_cli/test_update_host_obligation.py # tests/hermes_cli/test_update_import_guard.py # tests/hermes_cli/test_update_interrupted_recovery.py # tests/hermes_cli/test_update_inventory.py # tests/hermes_cli/test_update_launchd_unloaded_gateway.py # tests/hermes_cli/test_update_missing_configured_deps.py # tests/hermes_cli/test_update_modified_notice.py # tests/hermes_cli/test_update_multiplex_migration_hook.py # tests/hermes_cli/test_update_no_gateway_restart.py # tests/hermes_cli/test_update_orphan_backend_reap.py # tests/hermes_cli/test_update_parked_branch_guard.py # tests/hermes_cli/test_update_post_pull_syntax_guard.py # tests/hermes_cli/test_update_receipt.py # tests/hermes_cli/test_update_self_lock.py # tests/hermes_cli/test_update_shim_fail_closed.py # tests/hermes_cli/test_update_shim_self_lock.py # tests/hermes_cli/test_update_sqlite_remediation.py # tests/hermes_cli/test_update_stale_dashboard.py # tests/hermes_cli/test_update_stale_virtualenv.py # tests/hermes_cli/test_update_venv_health.py # tests/hermes_cli/test_update_venv_ownership_preflight.py # tests/hermes_cli/test_update_wedged_gateway.py # tests/hermes_cli/test_update_yes_flag.py # tests/hermes_cli/test_update_zip_two_phase.py # tests/hermes_cli/test_urllib_security.py # tests/hermes_cli/test_ux_messages_auth_config.py # tests/hermes_cli/test_ux_messages_startup.py # tests/hermes_cli/test_venv_holder_classifier.py # tests/hermes_cli/test_verify_console_scripts.py # tests/hermes_cli/test_verify_core_dependencies.py # tests/hermes_cli/test_web_server.py # tests/hermes_cli/test_web_server_console_ws.py # tests/hermes_cli/test_web_server_ws_ping.py # tests/hermes_cli/test_web_ui_build.py # tests/hermes_state/test_fts_rebuild_admission.py # tests/hermes_state/test_hermes_state.py # tests/plugins/memory/test_memory_lazy_install.py # tests/plugins/test_google_meet_plugin.py # tests/plugins/test_langfuse_plugin.py # tests/plugins/test_security_guidance_plugin.py # tests/plugins/test_transform_llm_output_hook.py # tests/scripts/desktop_update/test_desktop_update_windows_gateway_flag.py # tests/scripts/desktop_update/test_desktop_update_windows_python_handoff.py # tests/scripts/desktop_update/test_desktop_update_windows_timestamp.py # tests/scripts/install/test_install_clone_throttle_fallback.py # tests/scripts/install/test_install_lockfile_churn.py # tests/scripts/install/test_install_no_initial_commit.py # tests/scripts/install/test_install_sh_browser_install.py # tests/scripts/install/test_install_sh_node_prerelease.py # tests/scripts/install/test_install_sh_symlink_stomp.py # tests/scripts/install/test_install_sh_uv_lock_config.py # tests/scripts/install/test_install_unmerged_index.py # tests/scripts/test_contributor_map.py # tests/scripts/test_run_tests_parallel.py # tests/skills/test_competitor_news_monitor_skill.py # tests/skills/test_document_to_action_items_skill.py # tests/skills/test_google_workspace_setup.py # tests/skills/test_google_workspace_setup_deps.py # tests/skills/test_grounded_citations_skill.py # tests/skills/test_ip_as_logo_skill.py # tests/skills/test_live_dashboard_skill.py # tests/skills/test_mcp_oauth_remote_gateway_skill.py # tests/skills/test_office_document_skills.py # tests/skills/test_openclaw_migration.py # tests/skills/test_product_price_monitor_skill.py # tests/skills/test_scrollcraft_skill.py # tests/skills/test_setup_wizard_generator_skill.py # tests/skills/test_weekly_review_planning_skill.py # tests/test_engines_satisfiable.py # tests/test_fast_safe_load.py # tests/test_hermes_bootstrap.py # tests/test_hermes_constants.py # tests/test_hermes_logging.py # tests/test_managed_runtime_resolution.py # tests/test_model_tools_async_bridge.py # tests/test_packaging_build_guard.py # tests/test_packaging_metadata.py # tests/test_yaml_indent_consistency.py # tests/tools/test_approval_timeout_overflow.py # tests/tools/test_base_environment.py # tests/tools/test_bot_mode_dm.py # tests/tools/test_browser_chromium_check.py # tests/tools/test_browser_hardening.py # tests/tools/test_browser_homebrew_paths.py # tests/tools/test_browser_npx_warmup.py # tests/tools/test_browser_orphan_reaper.py # tests/tools/test_browser_real_profile.py # tests/tools/test_browser_use_cli.py # tests/tools/test_clipboard.py # tests/tools/test_code_execution.py # tests/tools/test_code_execution_modes.py # tests/tools/test_code_execution_windows_env.py # tests/tools/test_computer_use.py # tests/tools/test_delegate_liveness_timeout.py # tests/tools/test_execute_code_approval_cluster.py # tests/tools/test_execution_flag_detection.py # tests/tools/test_fal_common.py # tests/tools/test_file_operations.py # tests/tools/test_file_tools.py # tests/tools/test_file_tools_cwd_resolution.py # tests/tools/test_file_tools_live.py # tests/tools/test_lazy_deps.py # tests/tools/test_lazy_deps_durable_target.py # tests/tools/test_lazy_deps_managed.py # tests/tools/test_local_env_blocklist.py # tests/tools/test_local_tempdir.py # tests/tools/test_macos_protected_search.py # tests/tools/test_mcp_npx_cached_bin.py # tests/tools/test_oneshot_completion_linger.py # tests/tools/test_process_registry.py # tests/tools/test_read_file_schema_gating.py # tests/tools/test_skill_improvements.py # tests/tools/test_skills_sync.py # tests/tools/test_termux_api_detection.py # tests/tools/test_tirith_security.py # tests/tools/test_transcription_tools.py # tests/tools/test_tts_streaming.py # tests/tools/test_wake_word.py # tests/tui_gateway/test_compute_host_borrowed_lease.py # tests/tui_gateway/test_compute_host_turn_protocol.py # tests/tui_gateway/test_isolated_orphan_activity.py # tests/tui_gateway/test_protocol.py # tests/tui_gateway/test_slash_worker_profile_home.py # tests/tui_gateway/test_subprocess_encoding.py # tests/tui_gateway/test_tui_gateway_server.py # ui-tui/src/__tests__/terminalParity.test.ts # ui-tui/src/__tests__/termuxComposerLayout.test.ts # ui-tui/src/__tests__/textInputFastEcho.test.ts
415 lines
16 KiB
Python
415 lines
16 KiB
Python
"""Tests for the /review command engine — agent/review_engine.py.
|
|
|
|
Covers conversation snapshotting, reviewer-task composition,
|
|
auxiliary.review credential resolution, background dispatch through
|
|
delegate_task (including the internal per-call ``credentials_cfg``
|
|
override), and the shared dispatch-note formatter.
|
|
"""
|
|
|
|
import json
|
|
import time
|
|
from unittest.mock import MagicMock
|
|
|
|
import pytest
|
|
|
|
from agent import review_engine as re_mod
|
|
from agent.review_engine import (
|
|
build_review_task,
|
|
format_dispatch_note,
|
|
snapshot_recent_messages,
|
|
start_review,
|
|
)
|
|
from tools import async_delegation as ad
|
|
from tools.process_registry import process_registry
|
|
|
|
@pytest.fixture(autouse=True)
|
|
def _clean_state():
|
|
ad._reset_for_tests()
|
|
while not process_registry.completion_queue.empty():
|
|
process_registry.completion_queue.get_nowait()
|
|
yield
|
|
deadline = time.monotonic() + 2.0
|
|
while ad.active_count() and time.monotonic() < deadline:
|
|
time.sleep(0.02)
|
|
ad._reset_for_tests()
|
|
while not process_registry.completion_queue.empty():
|
|
process_registry.completion_queue.get_nowait()
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# snapshot_recent_messages
|
|
# ---------------------------------------------------------------------------
|
|
|
|
def test_snapshot_takes_last_ten_chat_messages_only():
|
|
msgs = (
|
|
[{"role": "system", "content": "sys"}]
|
|
+ [{"role": "user", "content": f"m{i}"} for i in range(15)]
|
|
+ [{"role": "tool", "content": "tool out"}]
|
|
)
|
|
snap = snapshot_recent_messages(msgs)
|
|
assert len(snap) == 10
|
|
assert snap[0]["text"] == "m5"
|
|
assert snap[-1]["text"] == "m14"
|
|
assert all(m["role"] == "user" for m in snap)
|
|
|
|
def test_snapshot_skips_toolcall_stub_assistant_messages():
|
|
msgs = [
|
|
{"role": "user", "content": "make a PR"},
|
|
{"role": "assistant", "content": "", "tool_calls": [{"id": "x"}]},
|
|
{"role": "tool", "content": "created"},
|
|
{"role": "assistant", "content": "PR #123: https://example.com/pr/123"},
|
|
]
|
|
snap = snapshot_recent_messages(msgs)
|
|
assert [m["text"] for m in snap] == [
|
|
"make a PR",
|
|
"PR #123: https://example.com/pr/123",
|
|
]
|
|
|
|
def test_snapshot_handles_multimodal_content_lists():
|
|
msgs = [{
|
|
"role": "user",
|
|
"content": [
|
|
{"type": "text", "text": "look at this"},
|
|
{"type": "image_url", "image_url": {"url": "x"}},
|
|
],
|
|
}]
|
|
snap = snapshot_recent_messages(msgs)
|
|
assert snap[0]["text"] == "look at this\n[image_url]"
|
|
|
|
def test_snapshot_caps_oversized_messages():
|
|
msgs = [{"role": "user", "content": "x" * 50_000}]
|
|
snap = snapshot_recent_messages(msgs)
|
|
assert len(snap[0]["text"]) < 13_000
|
|
assert snap[0]["text"].endswith("[... truncated ...]")
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# build_review_task
|
|
# ---------------------------------------------------------------------------
|
|
|
|
def test_build_review_task_includes_excerpt_and_prompt():
|
|
snap = [
|
|
{"role": "user", "text": "review my PR"},
|
|
{"role": "assistant", "text": "PR #99 opened"},
|
|
]
|
|
prompt = "focus on security\n" + "keep these instructions intact " * 20
|
|
goal, context = build_review_task(snap, prompt)
|
|
assert goal.startswith("Review: focus on security ")
|
|
assert len(goal) <= 80 and "\n" not in goal
|
|
assert goal.endswith("…")
|
|
assert re_mod._REVIEW_GOAL in context
|
|
assert "[USER]" in context and "[PRIMARY AGENT]" in context
|
|
assert "PR #99 opened" in context
|
|
assert prompt.strip() in context
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# auxiliary.review credential resolution
|
|
# ---------------------------------------------------------------------------
|
|
|
|
def test_load_review_credentials_cfg_reads_config(monkeypatch):
|
|
monkeypatch.setattr(
|
|
"hermes_cli.config.load_config_readonly",
|
|
lambda: {"auxiliary": {"review": {
|
|
"provider": "openrouter",
|
|
"model": "anthropic/claude-opus-4.6",
|
|
}}},
|
|
)
|
|
cfg = re_mod._load_review_credentials_cfg()
|
|
assert cfg == {
|
|
"provider": "openrouter",
|
|
"model": "anthropic/claude-opus-4.6",
|
|
"base_url": "",
|
|
"api_key": "",
|
|
"api_mode": "",
|
|
}
|
|
|
|
def test_load_review_credentials_cfg_auto_means_inherit(monkeypatch):
|
|
monkeypatch.setattr(
|
|
"hermes_cli.config.load_config_readonly",
|
|
lambda: {"auxiliary": {"review": {"provider": "auto", "model": ""}}},
|
|
)
|
|
assert re_mod._load_review_credentials_cfg() is None
|
|
|
|
def test_load_review_credentials_cfg_missing_section(monkeypatch):
|
|
monkeypatch.setattr(
|
|
"hermes_cli.config.load_config_readonly", lambda: {"auxiliary": {}}
|
|
)
|
|
assert re_mod._load_review_credentials_cfg() is None
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# delegate_task credentials_cfg override (the internal /review routing hook)
|
|
# ---------------------------------------------------------------------------
|
|
|
|
def _fake_parent():
|
|
parent = MagicMock()
|
|
parent._delegate_depth = 0
|
|
parent.session_id = "review-parent-sess"
|
|
parent._interrupt_requested = False
|
|
parent._active_children = []
|
|
parent._active_children_lock = None
|
|
return parent
|
|
|
|
def test_delegate_task_credentials_cfg_overrides_delegation_config(monkeypatch):
|
|
"""The per-call credentials_cfg dict must reach the credential resolver
|
|
instead of the global delegation config section."""
|
|
import tools.delegate_tool as dt
|
|
|
|
seen = {}
|
|
|
|
def fake_resolve(cfg, parent_agent):
|
|
seen["cfg"] = cfg
|
|
return {
|
|
"model": cfg.get("model"), "provider": None, "base_url": None,
|
|
"api_key": None, "api_mode": None, "command": None, "args": None,
|
|
}
|
|
|
|
fake_child = MagicMock()
|
|
fake_child._delegate_role = "leaf"
|
|
monkeypatch.setattr(dt, "_resolve_delegation_credentials", fake_resolve)
|
|
monkeypatch.setattr(dt, "_build_child_agent", lambda **kw: fake_child)
|
|
monkeypatch.setattr(
|
|
dt, "_run_single_child",
|
|
lambda *a, **k: {
|
|
"task_index": 0, "status": "completed", "summary": "ok",
|
|
"api_calls": 1, "duration_seconds": 0.1, "model": "m",
|
|
"exit_reason": "completed",
|
|
},
|
|
)
|
|
|
|
override = {"provider": "openrouter", "model": "review-model-x"}
|
|
out = dt.delegate_task(
|
|
goal="review this",
|
|
background=True,
|
|
parent_agent=_fake_parent(),
|
|
credentials_cfg=override,
|
|
)
|
|
parsed = json.loads(out)
|
|
assert parsed["status"] == "dispatched"
|
|
assert seen["cfg"] == override
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# start_review end-to-end through the async delegation rail
|
|
# ---------------------------------------------------------------------------
|
|
|
|
def test_start_review_dispatches_background_and_completes(monkeypatch):
|
|
import tools.delegate_tool as dt
|
|
|
|
captured = {}
|
|
|
|
def fake_run_single_child(task_index, goal, child=None, parent_agent=None, **kw):
|
|
captured["goal"] = goal
|
|
return {
|
|
"task_index": 0, "status": "completed",
|
|
"summary": "REVIEW: looks good", "api_calls": 2,
|
|
"duration_seconds": 0.1, "model": "m", "exit_reason": "completed",
|
|
}
|
|
|
|
fake_child = MagicMock()
|
|
fake_child._delegate_role = "leaf"
|
|
creds = {
|
|
"model": "m", "provider": None, "base_url": None, "api_key": None,
|
|
"api_mode": None, "command": None, "args": None,
|
|
}
|
|
built = {}
|
|
|
|
def fake_build(**kw):
|
|
built.update(kw)
|
|
return fake_child
|
|
|
|
monkeypatch.setattr(dt, "_build_child_agent", fake_build)
|
|
monkeypatch.setattr(dt, "_run_single_child", fake_run_single_child)
|
|
monkeypatch.setattr(dt, "_resolve_delegation_credentials", lambda *a, **k: creds)
|
|
monkeypatch.setattr(re_mod, "_load_review_credentials_cfg", lambda: None)
|
|
|
|
msgs = [
|
|
{"role": "user", "content": "open a PR for the fix"},
|
|
{"role": "assistant", "content": "PR #77 opened: https://x/pull/77"},
|
|
]
|
|
result = start_review(_fake_parent(), msgs, "check the tests")
|
|
assert result["status"] == "dispatched"
|
|
|
|
# The reviewer briefing carries the conversation excerpt + user prompt.
|
|
assert "PR #77 opened" in built["context"]
|
|
assert "check the tests" in built["context"]
|
|
assert built["goal"].startswith("Review: ")
|
|
assert re_mod._REVIEW_GOAL in built["context"]
|
|
|
|
# The completion re-enters via the shared queue like any subagent.
|
|
deadline = time.monotonic() + 5.0
|
|
evt = None
|
|
while time.monotonic() < deadline:
|
|
try:
|
|
evt = process_registry.completion_queue.get(timeout=0.2)
|
|
break
|
|
except Exception:
|
|
continue
|
|
assert evt is not None and evt["type"] == "async_delegation"
|
|
assert evt["results"][0]["summary"] == "REVIEW: looks good"
|
|
|
|
def test_start_review_rejects_empty_conversation():
|
|
with pytest.raises(ValueError, match="empty"):
|
|
start_review(_fake_parent(), [], "")
|
|
|
|
def test_start_review_requires_agent():
|
|
with pytest.raises(ValueError, match="No active agent"):
|
|
start_review(None, [{"role": "user", "content": "x"}], "")
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# collect_parent_loaded_skills — reviewer inherits the parent's working skills
|
|
# ---------------------------------------------------------------------------
|
|
|
|
def test_collect_skills_from_preloaded_prompt_and_history():
|
|
from agent.review_engine import collect_parent_loaded_skills
|
|
|
|
parent = MagicMock()
|
|
parent.ephemeral_system_prompt = (
|
|
'[IMPORTANT: The user launched this CLI session with the '
|
|
'"hermes-agent-dev" skill preloaded. Treat its instructions as '
|
|
'active guidance for the duration of this session unless the user '
|
|
'overrides them.]'
|
|
)
|
|
msgs = [
|
|
{"role": "assistant", "content": "", "tool_calls": [
|
|
{"function": {"name": "skill_view",
|
|
"arguments": '{"name": "github-pr-workflow"}'}},
|
|
# reference-file read of an already-counted skill: skipped
|
|
{"function": {"name": "skill_view",
|
|
"arguments": '{"name": "hermes-agent-dev", '
|
|
'"file_path": "references/x.md"}'}},
|
|
{"function": {"name": "read_file",
|
|
"arguments": '{"path": "/tmp/x"}'}},
|
|
]},
|
|
# duplicate load coalesces
|
|
{"role": "assistant", "content": "", "tool_calls": [
|
|
{"function": {"name": "skill_view",
|
|
"arguments": '{"name": "github-pr-workflow"}'}},
|
|
]},
|
|
]
|
|
names = collect_parent_loaded_skills(parent, msgs)
|
|
assert names == ["hermes-agent-dev", "github-pr-workflow"]
|
|
|
|
def test_collect_skills_empty_when_none_loaded():
|
|
from agent.review_engine import collect_parent_loaded_skills
|
|
|
|
parent = MagicMock()
|
|
parent.ephemeral_system_prompt = None
|
|
assert collect_parent_loaded_skills(parent, [
|
|
{"role": "user", "content": "hi"},
|
|
]) == []
|
|
|
|
def test_collect_skills_caps_at_limit():
|
|
from agent.review_engine import collect_parent_loaded_skills
|
|
|
|
parent = MagicMock()
|
|
parent.ephemeral_system_prompt = ""
|
|
msgs = [
|
|
{"role": "assistant", "content": "", "tool_calls": [
|
|
{"function": {"name": "skill_view",
|
|
"arguments": json.dumps({"name": f"skill-{i}"})}}
|
|
]}
|
|
for i in range(15)
|
|
]
|
|
import inspect
|
|
cap = inspect.signature(collect_parent_loaded_skills).parameters["limit"].default
|
|
assert len(collect_parent_loaded_skills(parent, msgs)) == cap < 15
|
|
|
|
def test_briefing_includes_loaded_skills_instruction():
|
|
snap = [{"role": "user", "text": "review my PR"}]
|
|
_, context = build_review_task(snap, "", ["hermes-agent-dev", "xitter"])
|
|
assert "hermes-agent-dev, xitter" in context
|
|
assert "skill_view" in context
|
|
|
|
def test_start_review_threads_loaded_skills_into_context(monkeypatch):
|
|
import tools.delegate_tool as dt
|
|
|
|
fake_child = MagicMock()
|
|
fake_child._delegate_role = "leaf"
|
|
creds = {
|
|
"model": "m", "provider": None, "base_url": None, "api_key": None,
|
|
"api_mode": None, "command": None, "args": None,
|
|
}
|
|
built = {}
|
|
|
|
def fake_build(**kw):
|
|
built.update(kw)
|
|
return fake_child
|
|
|
|
monkeypatch.setattr(dt, "_build_child_agent", fake_build)
|
|
monkeypatch.setattr(dt, "_resolve_delegation_credentials", lambda *a, **k: creds)
|
|
monkeypatch.setattr(
|
|
dt, "_run_single_child",
|
|
lambda *a, **k: {
|
|
"task_index": 0, "status": "completed", "summary": "ok",
|
|
"api_calls": 1, "duration_seconds": 0.1, "model": "m",
|
|
"exit_reason": "completed",
|
|
},
|
|
)
|
|
monkeypatch.setattr(re_mod, "_load_review_credentials_cfg", lambda: None)
|
|
|
|
parent = _fake_parent()
|
|
parent.ephemeral_system_prompt = (
|
|
'session with the "hermes-agent-dev" skill preloaded.'
|
|
)
|
|
msgs = [
|
|
{"role": "user", "content": "open a PR"},
|
|
{"role": "assistant", "content": "PR #9 opened"},
|
|
]
|
|
result = start_review(parent, msgs, "")
|
|
assert result["status"] == "dispatched"
|
|
assert "hermes-agent-dev" in built["context"]
|
|
assert "skill_view" in built["context"]
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# Workspace context files — ALL subagents (reviewer included) get AGENTS.md
|
|
# et al. in their child system prompt (tools/delegate_tool.py)
|
|
# ---------------------------------------------------------------------------
|
|
|
|
def test_child_system_prompt_embeds_workspace_context(tmp_path):
|
|
"""Real file I/O through the same loader the main system prompt uses."""
|
|
from tools.delegate_tool import _build_child_system_prompt
|
|
|
|
workspace = tmp_path / "proj"
|
|
workspace.mkdir()
|
|
(workspace / "AGENTS.md").write_text(
|
|
"# Project Rules\nAll fixes must cover sibling call sites.\n"
|
|
)
|
|
prompt = _build_child_system_prompt(
|
|
"do the thing", None, workspace_path=str(workspace)
|
|
)
|
|
assert "sibling call sites" in prompt
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# Registry sync: `review` must be a first-class slot in every aux-task surface
|
|
# ---------------------------------------------------------------------------
|
|
|
|
def test_review_registered_in_every_aux_surface():
|
|
"""The /review slot must appear in every aux-model picker registry.
|
|
|
|
Same contract as curator's registry test in tests/agent/test_curator.py:
|
|
DEFAULT_CONFIG schema, CLI picker (_AUX_TASKS), and dashboard REST
|
|
allowlist (_AUX_TASK_SLOTS). The desktop and web AUX_TASKS tsx arrays
|
|
mirror _AUX_TASK_SLOTS by convention (shared "Must match" comments).
|
|
"""
|
|
from hermes_cli.config import DEFAULT_CONFIG
|
|
from hermes_cli.main_provider_setup import _AUX_TASKS
|
|
from hermes_cli.web_server_config import _AUX_TASK_SLOTS
|
|
|
|
assert "review" in DEFAULT_CONFIG["auxiliary"], \
|
|
"review missing from DEFAULT_CONFIG['auxiliary']"
|
|
|
|
aux_keys = {k for k, _name, _desc in _AUX_TASKS}
|
|
assert "review" in aux_keys, "review missing from _AUX_TASKS (CLI picker)"
|
|
|
|
assert "review" in _AUX_TASK_SLOTS, \
|
|
"review missing from _AUX_TASK_SLOTS (dashboard REST API)"
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# format_dispatch_note
|
|
# ---------------------------------------------------------------------------
|
|
|
|
def test_format_dispatch_note_dispatched():
|
|
prompt = "security\n" * 100
|
|
note = format_dispatch_note(
|
|
{"status": "dispatched", "review_model": "opus"}, prompt
|
|
)
|
|
assert "security" not in note and "\n" not in note
|
|
assert len(note) < 80
|