Raise ruff line-length to 100 and extend the local pre-commit format pipeline (def-signature magic-comma normalization, short multi-line assert collapse, kwarg '=' spacing, blank-line-after-short-import removal, adjacent string-literal / f-string+plain merge, redundant-pass pruning). Every transform re-checks the file AST and is dropped if it would differ; the whole-repo reformat is verified AST-identical per file and idempotent.
182 lines
6.8 KiB
Python
182 lines
6.8 KiB
Python
# SPDX-License-Identifier: AGPL-3.0-only
|
|
# Copyright 2026-present the Unsloth AI Inc. team. All rights reserved. See /studio/LICENSE.AGPL-3.0
|
|
|
|
"""
|
|
Regression tests for the export log ring-buffer cursor semantics.
|
|
|
|
Context: the live export log SSE stream has a race where the frontend
|
|
opens the SSE connection AFTER the POST that starts the export. Any
|
|
lines the worker subprocess emits during the gap between POST and SSE
|
|
connect get buffered with seqs 1..k, and then the SSE default cursor
|
|
`get_current_log_seq()` returns k -- so lines 1..k are forever
|
|
unreachable to that client.
|
|
|
|
Fix: `clear_logs()` snapshots the pre-run seq into `_run_start_seq`
|
|
(exposed via `get_run_start_seq()`), and `routes/export.py` defaults
|
|
the SSE cursor to that snapshot instead of the current seq. Every line
|
|
appended during the current run has seq strictly greater than the
|
|
snapshot, so the client sees the full run regardless of when it
|
|
connects.
|
|
|
|
These tests exercise the orchestrator-side contract only (no
|
|
subprocess, no FastAPI, no frontend). The routes-level integration
|
|
with get_run_start_seq() is a one-line edit covered by manual testing
|
|
and the frontend build.
|
|
"""
|
|
|
|
from __future__ import annotations
|
|
|
|
import sys
|
|
import types
|
|
from pathlib import Path
|
|
|
|
import pytest
|
|
|
|
|
|
# Backend root on sys.path so `from core.export.orchestrator import ...`
|
|
# and friends resolve without the studio app bootstrap.
|
|
_BACKEND_DIR = Path(__file__).resolve().parent.parent
|
|
if str(_BACKEND_DIR) not in sys.path:
|
|
sys.path.insert(0, str(_BACKEND_DIR))
|
|
|
|
# ExportOrchestrator imports structlog and a few heavy modules at the
|
|
# top of orchestrator.py. Stub the ones we don't need in these unit
|
|
# tests so the import succeeds on machines without the full studio
|
|
# venv.
|
|
_loggers_stub = types.ModuleType("loggers")
|
|
_loggers_stub.get_logger = lambda name: __import__("logging").getLogger(name)
|
|
sys.modules.setdefault("loggers", _loggers_stub)
|
|
|
|
# structlog is only used for a module-level import; a bare stub is
|
|
# enough because we never call into it in these tests.
|
|
sys.modules.setdefault("structlog", types.ModuleType("structlog"))
|
|
|
|
# utils.paths.outputs_root is only called inside scan_checkpoints which
|
|
# we don't hit in these tests. Provide a stub module so the top-level
|
|
# import in orchestrator.py resolves.
|
|
_utils_pkg = types.ModuleType("utils")
|
|
_utils_pkg.__path__ = [] # mark as package
|
|
_utils_paths_stub = types.ModuleType("utils.paths")
|
|
_utils_paths_stub.outputs_root = lambda: Path("/tmp")
|
|
sys.modules.setdefault("utils", _utils_pkg)
|
|
sys.modules.setdefault("utils.paths", _utils_paths_stub)
|
|
|
|
|
|
@pytest.fixture
|
|
def orchestrator():
|
|
"""Fresh ExportOrchestrator with only the log-buffer state exercised."""
|
|
from core.export.orchestrator import ExportOrchestrator
|
|
return ExportOrchestrator()
|
|
|
|
|
|
def _append(
|
|
orch,
|
|
line: str,
|
|
stream: str = "stdout",
|
|
) -> None:
|
|
"""Shortcut for simulating a worker log message."""
|
|
orch._append_log({"type": "log", "stream": stream, "line": line, "ts": 0.0})
|
|
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# clear_logs() semantics
|
|
# ---------------------------------------------------------------------------
|
|
|
|
|
|
def test_run_start_seq_is_zero_before_any_logs(orchestrator) -> None:
|
|
"""A brand-new orchestrator must report run_start_seq == 0 so a
|
|
first SSE connection picks up every line from seq 1 onward."""
|
|
assert orchestrator.get_run_start_seq() == 0
|
|
|
|
|
|
def test_clear_logs_snapshots_current_seq(orchestrator) -> None:
|
|
"""clear_logs() must capture _log_seq BEFORE clearing the buffer,
|
|
so subsequent runs can anchor their SSE cursor at the snapshot."""
|
|
_append(orchestrator, "old run line 1")
|
|
_append(orchestrator, "old run line 2")
|
|
_append(orchestrator, "old run line 3")
|
|
assert orchestrator.get_current_log_seq() == 3
|
|
|
|
orchestrator.clear_logs()
|
|
|
|
assert orchestrator.get_run_start_seq() == 3
|
|
assert orchestrator.get_current_log_seq() == 3 # seq counter preserved
|
|
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# Race regression: SSE connects AFTER lines have been emitted
|
|
# ---------------------------------------------------------------------------
|
|
|
|
|
|
def test_sse_default_cursor_catches_all_current_run_lines(orchestrator) -> None:
|
|
"""Simulate the POST-then-SSE race: worker starts emitting lines
|
|
immediately after clear_logs(), SSE connects several lines later.
|
|
Using get_run_start_seq() as the default cursor MUST return every
|
|
line emitted since clear_logs() ran.
|
|
|
|
Pre-fix, the SSE defaulted to get_current_log_seq() at connect
|
|
time, which would return the last-seen seq and miss lines N+1..M.
|
|
"""
|
|
# Previous run leaves some buffered lines.
|
|
_append(orchestrator, "previous run line A")
|
|
_append(orchestrator, "previous run line B")
|
|
|
|
# New run starts: orchestrator clears the buffer and snapshots seq.
|
|
orchestrator.clear_logs()
|
|
run_start = orchestrator.get_run_start_seq()
|
|
|
|
# Worker emits early lines BEFORE the SSE connects.
|
|
_append(orchestrator, "Importing Unsloth...")
|
|
_append(orchestrator, "Loading checkpoint: /foo/bar")
|
|
_append(orchestrator, "Starting export...")
|
|
|
|
# SSE connects now and asks "give me everything after the run
|
|
# start cursor".
|
|
entries, new_cursor = orchestrator.get_logs_since(run_start)
|
|
|
|
# All three early lines must be present. Pre-fix this was [].
|
|
lines = [e["line"] for e in entries]
|
|
assert lines == [
|
|
"Importing Unsloth...",
|
|
"Loading checkpoint: /foo/bar",
|
|
"Starting export...",
|
|
]
|
|
assert new_cursor == entries[-1]["seq"]
|
|
|
|
|
|
def test_sse_default_cursor_excludes_previous_run(orchestrator) -> None:
|
|
"""After clear_logs(), lines from the PREVIOUS run must not leak
|
|
into the new run's SSE stream. Pre-fix this worked correctly
|
|
(clear_logs cleared the deque); the fix must preserve it.
|
|
"""
|
|
_append(orchestrator, "previous run line 1")
|
|
_append(orchestrator, "previous run line 2")
|
|
_append(orchestrator, "previous run line 3")
|
|
assert orchestrator.get_current_log_seq() == 3
|
|
|
|
orchestrator.clear_logs()
|
|
run_start = orchestrator.get_run_start_seq()
|
|
|
|
_append(orchestrator, "new run line")
|
|
|
|
entries, _ = orchestrator.get_logs_since(run_start)
|
|
assert [e["line"] for e in entries] == ["new run line"]
|
|
|
|
|
|
def test_clear_logs_twice_advances_run_start(orchestrator) -> None:
|
|
"""Back-to-back clear_logs() calls (e.g. cleanup -> load ->
|
|
export in the same dialog session) must each re-anchor run_start
|
|
at the current seq, so successive runs each start with a fresh
|
|
low-water mark."""
|
|
_append(orchestrator, "run 1 line a")
|
|
_append(orchestrator, "run 1 line b")
|
|
|
|
orchestrator.clear_logs()
|
|
assert orchestrator.get_run_start_seq() == 2
|
|
|
|
_append(orchestrator, "run 2 line a")
|
|
_append(orchestrator, "run 2 line b")
|
|
_append(orchestrator, "run 2 line c")
|
|
|
|
orchestrator.clear_logs()
|
|
assert orchestrator.get_run_start_seq() == 5
|