Files
Thuc Pham 3960618454 chore: create-llama monorepo (#581)
* chore: create-llama monorepo

* add root package.json and pnpm workspace

* keep e2e inside create-llama

* update root package.json

* move scripts and dev dependencies of create-llama to root

* update e2e test for create-llama package

* update lint workflow

* update release llama-index-server workflow

* update path for test_llama_index_server workflow

* remove local lock file

* keep lint and format in create-llama

* fix: format

* update pre-commit

* move playwright back to create-llama

* disable pnpm for installing generated frontend

* use npm for type check

* update gitignore

* try --ignore-workspace option

* Move llama-index-server from packages/python-server to python directory

* update CI for python server

* Create plenty-spies-tickle.md
2025-04-25 18:38:02 +07:00

150 lines
5.1 KiB
Python

import logging
from unittest.mock import AsyncMock, MagicMock
import pytest
from fastapi import FastAPI
from httpx import ASGITransport, AsyncClient
from llama_index.core.workflow import StopEvent, Workflow
from llama_index.core.workflow.handler import WorkflowHandler
from llama_index.server.api.models import ChatAPIMessage, ChatRequest
from llama_index.server.api.routers.chat import chat_router
@pytest.fixture()
def logger():
return logging.getLogger("test")
@pytest.fixture()
def chat_request():
"""Create a simple chat request with one user message."""
return ChatRequest(
messages=[ChatAPIMessage(role="user", content="Hello, how are you?")]
)
@pytest.fixture()
def mock_workflow():
"""Create a mock workflow that returns a simple response."""
workflow = MagicMock(spec=Workflow)
handler = AsyncMock(spec=WorkflowHandler)
# Setup the handler to stream a simple response event
async def mock_stream_events():
yield StopEvent(result="I'm doing well, thank you for asking!")
handler.stream_events.return_value = mock_stream_events()
workflow.run.return_value = handler
return workflow
@pytest.fixture()
def workflow_factory(mock_workflow):
"""Create a factory function that returns our mock workflow."""
def factory(verbose=False):
return mock_workflow
return factory
@pytest.mark.asyncio()
async def test_chat_router(chat_request, workflow_factory, logger):
"""Test that the chat router handles a request correctly."""
# Create a FastAPI app and mount our router
app = FastAPI()
router = chat_router(workflow_factory, logger)
app.include_router(router)
# Make a request to the chat endpoint
async with AsyncClient(
transport=ASGITransport(app=app), base_url="http://test"
) as client:
response = await client.post("/chat", json=chat_request.model_dump())
# Check response status
assert response.status_code == 200
# For streaming responses we don't check the content-type header directly
# Instead, check that we get the expected content in the response body
# The response is a stream, so we need to collect the chunks
content = response.content.decode()
# Verify content structure follows expected format
assert "0:" in content # Text prefix for VercelStreamResponse
# Verify if the response contains the expected message
assert "I'm doing well" in content
# Verify the mock workflow was called correctly
mock_workflow = workflow_factory()
mock_workflow.run.assert_called_once()
# Verify the workflow was called with the correct arguments
call_args = mock_workflow.run.call_args[1]
assert call_args["user_msg"] == "Hello, how are you?"
assert isinstance(call_args["chat_history"], list)
assert len(call_args["chat_history"]) == 0 # No history for first message
@pytest.mark.asyncio()
async def test_chat_with_agent_workflow(logger):
"""Test that the chat router works with a workflow that mimics an agent workflow."""
# Create a simple workflow that mimics an agent workflow
mock_workflow = MagicMock(spec=Workflow)
handler = AsyncMock(spec=WorkflowHandler)
# Setup the handler to stream a simple response about weather
async def mock_stream_events():
yield StopEvent(
result="The weather in New York is sunny. I used the weather tool to get this information."
)
handler.stream_events.return_value = mock_stream_events()
mock_workflow.run.return_value = handler
# Create a factory function that returns our mock workflow
def workflow_factory(verbose=False):
return mock_workflow
# Create a FastAPI app and mount our router
app = FastAPI()
router = chat_router(workflow_factory, logger)
app.include_router(router)
# Create a chat request asking about weather
chat_request = ChatRequest(
messages=[
ChatAPIMessage(role="user", content="What's the weather in New York?")
]
)
# Make a request to the chat endpoint
async with AsyncClient(
transport=ASGITransport(app=app), base_url="http://test"
) as client:
response = await client.post("/chat", json=chat_request.model_dump())
# Check response status
assert response.status_code == 200
# The response is a stream, so we need to collect the chunks
content = response.content.decode()
# Verify content structure follows expected format
assert "0:" in content # Text prefix for VercelStreamResponse
# Verify the response content contains expected keywords
assert "weather" in content and "New York" in content and "sunny" in content
# Verify the mock workflow was called correctly
mock_workflow.run.assert_called_once()
# Verify the workflow was called with the correct arguments
call_args = mock_workflow.run.call_args[1]
assert call_args["user_msg"] == "What's the weather in New York?"
assert isinstance(call_args["chat_history"], list)
assert len(call_args["chat_history"]) == 0 # No history for first message