mirror of
https://github.com/NousResearch/hermes-agent.git
synced 2026-07-24 16:54:43 +00:00
The un-mentioned wake decision relied on three checks: thread root in _bot_message_ts (populated only by the adapter's own send() path), _mentioned_threads (populated on @mention), and an existing session. Two gaps (#63530): - Gap A: bot messages posted OUTSIDE gateway send() — skills/scripts calling chat.postMessage directly, cron/API posts — never enter _bot_message_ts, so human replies in those threads were silently dropped. - Gap B: _bot_message_ts is process memory; after a gateway restart the bot stopped waking on replies to threads it started before the restart. Fix: add a 4th, API-derived check — _bot_authored_thread_root — which resolves the thread root's author via conversations.replies (cached in _thread_context_cache via the new parent_user_id field, TTL-bounded). Root authorship comes from Slack itself, so it covers outside-send posts and survives restarts, unlike in-memory ts tracking. The wake decision is extracted into _should_wake_on_unmentioned_message for direct unit testing; the legacy checks remain first (cheap, additive). Fixes #63530. Salvaged from #64067 by @knoal (author metadata normalized to their GitHub identity), rebased over the thread-context formatter split and extended with per-team bot-id resolution for multi-workspace installs.
329 lines
12 KiB
Python
329 lines
12 KiB
Python
"""Regression tests for #63530 — Slack adapter drops human replies in
|
|
threads whose root was posted by the bot via direct chat.postMessage
|
|
(outside the gateway's send() path).
|
|
|
|
Background: the adapter's wake-decision at the un-mentioned branch in
|
|
_handle_slack_message uses three checks:
|
|
|
|
1. thread_ts ∈ _bot_message_ts (only populated by send() / files_upload_v2)
|
|
2. thread_ts ∈ _mentioned_threads (only populated on @mention)
|
|
3. _has_active_session_for_thread(...) (survives restarts)
|
|
|
|
When a skill posts a triage message into a Slack thread via the Web API
|
|
directly (chat.postMessage, no gateway run), the bot's own ts is NOT
|
|
recorded in _bot_message_ts. A human reply in that thread, without an
|
|
@-mention and without an existing session, falls through all three
|
|
checks and is silently dropped. The same gap opens after a gateway
|
|
restart: _bot_message_ts is process memory, so threads the bot started
|
|
before the restart no longer wake it.
|
|
|
|
Fix: a 4th check — was the thread root authored by the bot? Root
|
|
authorship is derived from the Slack API (conversations.replies), so it
|
|
survives restarts, unlike the in-memory ts set. The wake decision is
|
|
extracted into _should_wake_on_unmentioned_message so it's directly
|
|
testable without spinning up Slack.
|
|
"""
|
|
|
|
import sys
|
|
from unittest.mock import AsyncMock, MagicMock
|
|
|
|
import pytest
|
|
|
|
|
|
# Mock slack-bolt / slack-sdk the same way test_slack_mention.py does.
|
|
def _ensure_slack_mock():
|
|
if "slack_bolt" in sys.modules and hasattr(sys.modules["slack_bolt"], "__file__"):
|
|
return
|
|
slack_bolt = MagicMock()
|
|
slack_bolt.async_app.AsyncApp = MagicMock
|
|
slack_bolt.adapter.socket_mode.async_handler.AsyncSocketModeHandler = MagicMock
|
|
slack_sdk = MagicMock()
|
|
slack_sdk.web.async_client.AsyncWebClient = MagicMock
|
|
for name, mod in [
|
|
("slack_bolt", slack_bolt),
|
|
("slack_bolt.async_app", slack_bolt.async_app),
|
|
("slack_bolt.adapter", slack_bolt.adapter),
|
|
("slack_bolt.adapter.socket_mode", slack_bolt.adapter.socket_mode),
|
|
(
|
|
"slack_bolt.adapter.socket_mode.async_handler",
|
|
slack_bolt.adapter.socket_mode.async_handler,
|
|
),
|
|
("slack_sdk", slack_sdk),
|
|
("slack_sdk.web", slack_sdk.web),
|
|
("slack_sdk.web.async_client", slack_sdk.web.async_client),
|
|
]:
|
|
sys.modules.setdefault(name, mod)
|
|
sys.modules.setdefault("aiohttp", MagicMock())
|
|
|
|
|
|
_ensure_slack_mock()
|
|
|
|
import plugins.platforms.slack.adapter as _slack_mod # noqa: E402
|
|
|
|
_slack_mod.SLACK_AVAILABLE = True
|
|
|
|
from plugins.platforms.slack.adapter import ( # noqa: E402
|
|
SlackAdapter,
|
|
_ThreadContextCache,
|
|
)
|
|
|
|
from gateway.config import Platform, PlatformConfig # noqa: E402
|
|
|
|
|
|
BOT_USER_ID = "U_BOT_OWN"
|
|
CHANNEL_ID = "C_incident"
|
|
USER_ID = "U_engineer"
|
|
THREAD_TS = "1700000000.000100"
|
|
|
|
|
|
def _make_adapter(bot_authored_root: bool = False):
|
|
"""Build a bare SlackAdapter with the wake-decision state controlled.
|
|
|
|
None of the 3 legacy in-memory checks pass by default: the bot didn't
|
|
send via gateway, the thread wasn't @-mentioned, and there is no active
|
|
session — exactly the post-restart / outside-send state.
|
|
"""
|
|
adapter = object.__new__(SlackAdapter)
|
|
adapter.platform = Platform.SLACK
|
|
adapter.config = PlatformConfig(
|
|
enabled=True,
|
|
extra={"require_mention": True, "strict_mention": False},
|
|
)
|
|
adapter._bot_user_id = BOT_USER_ID
|
|
adapter._team_bot_user_ids = {}
|
|
adapter._bot_message_ts = set()
|
|
adapter._mentioned_threads = set()
|
|
adapter._thread_context_cache = {}
|
|
|
|
adapter._has_active_session_for_thread = lambda **kw: False
|
|
# Mock _fetch_thread_context so the miss-path doesn't make a real
|
|
# Slack API call. Tests that need a populated cache pre-populate
|
|
# _thread_context_cache directly.
|
|
adapter._fetch_thread_context = AsyncMock(return_value="")
|
|
# The 4th-check helper is mocked so wake-decision tests can control
|
|
# its result without setting up the full cache path. Helper-specific
|
|
# tests call the real method via the class instead.
|
|
adapter._bot_authored_thread_root = AsyncMock(return_value=bot_authored_root)
|
|
|
|
return adapter
|
|
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# _should_wake_on_unmentioned_message — composes all 4 checks
|
|
# ---------------------------------------------------------------------------
|
|
|
|
|
|
@pytest.mark.asyncio
|
|
async def test_wake_decision_returns_false_when_not_thread_reply():
|
|
"""A top-level channel message (no thread_ts) should never wake the bot
|
|
when require_mention is true — unchanged by this fix."""
|
|
adapter = _make_adapter(bot_authored_root=True)
|
|
wake = await adapter._should_wake_on_unmentioned_message(
|
|
event_thread_ts=None,
|
|
channel_id=CHANNEL_ID,
|
|
user_id=USER_ID,
|
|
is_thread_reply=False,
|
|
)
|
|
assert wake is False
|
|
|
|
|
|
@pytest.mark.asyncio
|
|
async def test_wake_decision_returns_false_when_all_four_checks_miss():
|
|
"""All four checks miss (no bot-message, no mention, no session, no
|
|
bot-authored root) → wake decision is False."""
|
|
adapter = _make_adapter(bot_authored_root=False)
|
|
wake = await adapter._should_wake_on_unmentioned_message(
|
|
event_thread_ts=THREAD_TS,
|
|
channel_id=CHANNEL_ID,
|
|
user_id=USER_ID,
|
|
is_thread_reply=True,
|
|
)
|
|
assert wake is False
|
|
|
|
|
|
@pytest.mark.asyncio
|
|
async def test_wake_decision_returns_true_when_bot_authored_thread_root():
|
|
"""The new behavior (#63530): a human reply in a thread whose root was
|
|
authored by the bot via direct chat.postMessage (outside gateway send)
|
|
wakes the bot even though none of the legacy 3 checks pass — including
|
|
after a restart, when _bot_message_ts is empty."""
|
|
adapter = _make_adapter(bot_authored_root=True)
|
|
wake = await adapter._should_wake_on_unmentioned_message(
|
|
event_thread_ts=THREAD_TS,
|
|
channel_id=CHANNEL_ID,
|
|
user_id=USER_ID,
|
|
is_thread_reply=True,
|
|
)
|
|
assert wake is True, (
|
|
"human reply in a thread whose root was bot-posted (not via gateway "
|
|
"send) should wake the bot — #63530"
|
|
)
|
|
|
|
|
|
@pytest.mark.asyncio
|
|
async def test_wake_decision_returns_true_when_legacy_check_1_hits():
|
|
"""Regression guard: _bot_message_ts hit still wakes (additive check)."""
|
|
adapter = _make_adapter(bot_authored_root=False)
|
|
adapter._bot_message_ts = {THREAD_TS}
|
|
wake = await adapter._should_wake_on_unmentioned_message(
|
|
event_thread_ts=THREAD_TS,
|
|
channel_id=CHANNEL_ID,
|
|
user_id=USER_ID,
|
|
is_thread_reply=True,
|
|
)
|
|
assert wake is True
|
|
|
|
|
|
@pytest.mark.asyncio
|
|
async def test_wake_decision_returns_true_when_legacy_check_2_hits():
|
|
"""Regression guard: _mentioned_threads hit still wakes (additive)."""
|
|
adapter = _make_adapter(bot_authored_root=False)
|
|
adapter._mentioned_threads = {THREAD_TS}
|
|
wake = await adapter._should_wake_on_unmentioned_message(
|
|
event_thread_ts=THREAD_TS,
|
|
channel_id=CHANNEL_ID,
|
|
user_id=USER_ID,
|
|
is_thread_reply=True,
|
|
)
|
|
assert wake is True
|
|
|
|
|
|
@pytest.mark.asyncio
|
|
async def test_wake_decision_returns_true_when_legacy_check_3_hits():
|
|
"""Regression guard: an active session still wakes (additive)."""
|
|
adapter = _make_adapter(bot_authored_root=False)
|
|
adapter._has_active_session_for_thread = lambda **kw: True
|
|
wake = await adapter._should_wake_on_unmentioned_message(
|
|
event_thread_ts=THREAD_TS,
|
|
channel_id=CHANNEL_ID,
|
|
user_id=USER_ID,
|
|
is_thread_reply=True,
|
|
)
|
|
assert wake is True
|
|
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# _bot_authored_thread_root — the API-derived, restart-surviving check
|
|
# ---------------------------------------------------------------------------
|
|
|
|
|
|
@pytest.mark.asyncio
|
|
async def test_bot_authored_thread_root_true_from_cache():
|
|
"""Cache hit whose parent_user_id matches the bot's user_id → True."""
|
|
adapter = _make_adapter()
|
|
adapter._thread_context_cache = {
|
|
f"{CHANNEL_ID}:{THREAD_TS}:": _ThreadContextCache(
|
|
content="[Thread context — prior messages...]",
|
|
fetched_at=0,
|
|
message_count=1,
|
|
parent_text="triage analysis",
|
|
parent_user_id=BOT_USER_ID,
|
|
),
|
|
}
|
|
|
|
result = await SlackAdapter._bot_authored_thread_root(
|
|
adapter, CHANNEL_ID, THREAD_TS
|
|
)
|
|
assert result is True
|
|
|
|
|
|
@pytest.mark.asyncio
|
|
async def test_bot_authored_thread_root_false_for_human_authored_root():
|
|
"""A human-authored root must return False even on a cache hit — guards
|
|
against waking on any thread reply just because the cache is warm."""
|
|
adapter = _make_adapter()
|
|
adapter._thread_context_cache = {
|
|
f"{CHANNEL_ID}:{THREAD_TS}:": _ThreadContextCache(
|
|
content="[Thread context — prior messages...]",
|
|
fetched_at=0,
|
|
message_count=1,
|
|
parent_text="someone else's message",
|
|
parent_user_id="U_other_user",
|
|
),
|
|
}
|
|
|
|
result = await SlackAdapter._bot_authored_thread_root(
|
|
adapter, CHANNEL_ID, THREAD_TS
|
|
)
|
|
assert result is False
|
|
|
|
|
|
@pytest.mark.asyncio
|
|
async def test_bot_authored_thread_root_false_on_empty_thread_ts():
|
|
"""Defensive: empty thread_ts short-circuits to False without any
|
|
cache lookup or network call."""
|
|
adapter = _make_adapter()
|
|
result = await SlackAdapter._bot_authored_thread_root(adapter, CHANNEL_ID, "")
|
|
assert result is False
|
|
adapter._fetch_thread_context.assert_not_awaited()
|
|
|
|
|
|
@pytest.mark.asyncio
|
|
async def test_bot_authored_thread_root_fetches_on_cache_miss():
|
|
"""Cache miss → _fetch_thread_context runs; a successful fetch that
|
|
populates parent_user_id with the bot's id yields True. This is the
|
|
restart path: fresh process, empty caches, root authorship recovered
|
|
from the Slack API."""
|
|
adapter = _make_adapter()
|
|
|
|
async def _fake_fetch(channel_id, thread_ts, current_ts, team_id=""):
|
|
adapter._thread_context_cache[f"{channel_id}:{thread_ts}:{team_id}"] = (
|
|
_ThreadContextCache(
|
|
content="ctx",
|
|
fetched_at=0,
|
|
message_count=1,
|
|
parent_text="bot-posted root",
|
|
parent_user_id=BOT_USER_ID,
|
|
)
|
|
)
|
|
return "ctx"
|
|
|
|
adapter._fetch_thread_context = AsyncMock(side_effect=_fake_fetch)
|
|
|
|
result = await SlackAdapter._bot_authored_thread_root(
|
|
adapter, CHANNEL_ID, THREAD_TS
|
|
)
|
|
assert result is True
|
|
adapter._fetch_thread_context.assert_awaited_once()
|
|
|
|
|
|
@pytest.mark.asyncio
|
|
async def test_bot_authored_thread_root_false_when_fetch_fails():
|
|
"""Fetch failure (empty result, nothing cached) → False, no wake."""
|
|
adapter = _make_adapter()
|
|
adapter._fetch_thread_context = AsyncMock(return_value="")
|
|
|
|
result = await SlackAdapter._bot_authored_thread_root(
|
|
adapter, CHANNEL_ID, THREAD_TS
|
|
)
|
|
assert result is False
|
|
|
|
|
|
@pytest.mark.asyncio
|
|
async def test_bot_authored_thread_root_uses_per_team_bot_id():
|
|
"""Multi-workspace: the comparison must use the team's bot user id,
|
|
not the primary workspace's."""
|
|
adapter = _make_adapter()
|
|
adapter._team_bot_user_ids = {"T2": "U_BOT_T2"}
|
|
adapter._thread_context_cache = {
|
|
f"{CHANNEL_ID}:{THREAD_TS}:T2": _ThreadContextCache(
|
|
content="ctx",
|
|
fetched_at=0,
|
|
message_count=1,
|
|
parent_text="root",
|
|
parent_user_id="U_BOT_T2",
|
|
),
|
|
}
|
|
|
|
result = await SlackAdapter._bot_authored_thread_root(
|
|
adapter, CHANNEL_ID, THREAD_TS, team_id="T2"
|
|
)
|
|
assert result is True
|
|
# And the primary bot id must NOT match in that workspace.
|
|
adapter._thread_context_cache[f"{CHANNEL_ID}:{THREAD_TS}:T2"].parent_user_id = (
|
|
BOT_USER_ID
|
|
)
|
|
result = await SlackAdapter._bot_authored_thread_root(
|
|
adapter, CHANNEL_ID, THREAD_TS, team_id="T2"
|
|
)
|
|
assert result is False
|