Merged
Show file tree
Hide file tree
Changes from all commits
Commits
File filter

Filter by extension

Filter by extension

Conversations
Failed to load comments.
Loading
Jump to
Jump to file
Failed to load files.
Loading
Diff view
Diff view
2 changes: 1 addition & 1 deletion py/noxfile.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -259,7 +259,7 @@ def test_litellm(session, version):
# Install fastapi and orjson as they're required by litellm for proxy/responses operations
session.install("openai<=1.99.9", "--force-reinstall", "fastapi", "orjson")
_install(session, "litellm", version)
_run_tests(session, f"{WRAPPER_DIR}/test_litellm.py")
_run_tests(session, f"{INTEGRATION_DIR}/litellm/test_litellm.py")
_run_core_tests(session)


Expand Down
6 changes: 3 additions & 3 deletions py/src/braintrust/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -73,6 +73,9 @@ def is_equal(expected, output):
from .integrations.anthropic import (
wrap_anthropic, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.openrouter import (
wrap_openrouter, # noqa: F401 # type: ignore[reportUnusedImport]
)
Expand All@@ -98,6 +101,3 @@ def is_equal(expected, output):
BT_IS_ASYNC_ATTRIBUTE, # noqa: F401 # type: ignore[reportUnusedImport]
MarkAsyncWrapper, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .wrappers.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
11 changes: 2 additions & 9 deletions py/src/braintrust/auto.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -15,6 +15,7 @@
ClaudeAgentSDKIntegration,
DSPyIntegration,
GoogleGenAIIntegration,
LiteLLMIntegration,
OpenRouterIntegration,
PydanticAIIntegration,
)
Expand DownExpand Up@@ -123,7 +124,7 @@ def auto_instrument(
if anthropic:
results["anthropic"] = _instrument_integration(AnthropicIntegration)
if litellm:
results["litellm"] = _instrument_litellm()
results["litellm"] = _instrument_integration(LiteLLMIntegration)
if pydantic_ai:
results["pydantic_ai"] = _instrument_integration(PydanticAIIntegration)
if google_genai:
Expand DownExpand Up@@ -156,11 +157,3 @@ def _instrument_integration(integration) -> bool:
with _try_patch():
return integration.setup()
return False


def _instrument_litellm() -> bool:
with _try_patch():
from braintrust.wrappers.litellm import patch_litellm

return patch_litellm()
return False
2 changes: 2 additions & 0 deletions py/src/braintrust/integrations/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -5,6 +5,7 @@
from .claude_agent_sdk import ClaudeAgentSDKIntegration
from .dspy import DSPyIntegration
from .google_genai import GoogleGenAIIntegration
from .litellm import LiteLLMIntegration
from .openrouter import OpenRouterIntegration
from .pydantic_ai import PydanticAIIntegration

Expand All@@ -17,6 +18,7 @@
"ClaudeAgentSDKIntegration",
"DSPyIntegration",
"GoogleGenAIIntegration",
"LiteLLMIntegration",
"OpenRouterIntegration",
"PydanticAIIntegration",
]
Original file line numberDiff line numberDiff line change
@@ -1,24 +1,29 @@
"""Test auto_instrument for LiteLLM."""

from pathlib import Path

import litellm
from braintrust.auto import auto_instrument
from braintrust.integrations.litellm import LiteLLMIntegration
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

# 1. Verify not patched initially
assert not hasattr(litellm, "_braintrust_wrapped")
assert not LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 2. Instrument
results = auto_instrument()
assert results.get("litellm") == True
assert hasattr(litellm, "_braintrust_wrapped")
assert LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 3. Idempotent
results2 = auto_instrument()
assert results2.get("litellm") == True

# 4. Make API call and verify span
with autoinstrument_test_context("test_auto_litellm") as memory_logger:
with autoinstrument_test_context("test_auto_litellm", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Say hi"}],
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,17 +1,20 @@
"""Test that patch_litellm() patches aresponses."""

import asyncio
from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()


async def main():
with autoinstrument_test_context("test_patch_litellm_aresponses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_aresponses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = await litellm.aresponses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,13 +1,17 @@
"""Test that patch_litellm() patches responses."""

from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()

with autoinstrument_test_context("test_patch_litellm_responses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_responses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.responses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
2 changes: 1 addition & 1 deletion py/src/braintrust/integrations/dspy/tracing.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -50,7 +50,7 @@ class BraintrustDSpyCallback(BaseCallback):
and disable DSPy's disk cache:

```python
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
patch_litellm()

import dspy
Expand Down
40 changes: 40 additions & 0 deletions py/src/braintrust/integrations/litellm/__init__.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,40 @@
"""Braintrust LiteLLM integration."""

from .integration import LiteLLMIntegration
from .patchers import wrap_litellm


def patch_litellm() -> bool:
"""Patch LiteLLM to add Braintrust tracing.

This wraps litellm.completion, litellm.acompletion, litellm.responses,
litellm.aresponses, litellm.embedding, and litellm.moderation to
automatically create Braintrust spans with detailed token metrics,
timing, and costs.

Returns:
True if LiteLLM was patched (or already patched), False if LiteLLM is not installed.

Example:
```python
import braintrust
braintrust.integrations.litellm.patch_litellm()

import litellm
from braintrust import init_logger

logger = init_logger(project="my-project")
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}]
)
```
"""
return LiteLLMIntegration.setup()


__all__ = [
"LiteLLMIntegration",
"patch_litellm",
"wrap_litellm",
]
13 changes: 13 additions & 0 deletions py/src/braintrust/integrations/litellm/integration.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,13 @@
"""LiteLLM integration definition."""

from braintrust.integrations.base import BaseIntegration

from .patchers import _ALL_LITELLM_PATCHERS


class LiteLLMIntegration(BaseIntegration):
"""Braintrust instrumentation for the LiteLLM Python SDK."""

name = "litellm"
import_names = ("litellm",)
patchers = _ALL_LITELLM_PATCHERS
109 changes: 109 additions & 0 deletions py/src/braintrust/integrations/litellm/patchers.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,109 @@
"""LiteLLM patchers — FunctionWrapperPatcher subclasses for each patch target."""

from typing import Any

from braintrust.integrations.base import FunctionWrapperPatcher

from .tracing import (
_acompletion_wrapper_async,
_aresponses_wrapper_async,
_completion_wrapper,
_embedding_wrapper,
_moderation_wrapper,
_responses_wrapper,
)


# ---------------------------------------------------------------------------
# Individual patchers
# ---------------------------------------------------------------------------


class LiteLLMCompletionPatcher(FunctionWrapperPatcher):
name = "litellm.completion"
target_path = "completion"
wrapper = _completion_wrapper


class LiteLLMAcompletionPatcher(FunctionWrapperPatcher):
name = "litellm.acompletion"
target_path = "acompletion"
wrapper = _acompletion_wrapper_async


class LiteLLMResponsesPatcher(FunctionWrapperPatcher):
name = "litellm.responses"
target_path = "responses"
wrapper = _responses_wrapper


class LiteLLMAresponsesPatcher(FunctionWrapperPatcher):
name = "litellm.aresponses"
target_path = "aresponses"
wrapper = _aresponses_wrapper_async


class LiteLLMEmbeddingPatcher(FunctionWrapperPatcher):
name = "litellm.embedding"
target_path = "embedding"
wrapper = _embedding_wrapper


class LiteLLMModerationPatcher(FunctionWrapperPatcher):
name = "litellm.moderation"
target_path = "moderation"
wrapper = _moderation_wrapper


# ---------------------------------------------------------------------------
# All patchers, in declaration order
# ---------------------------------------------------------------------------

_ALL_LITELLM_PATCHERS = (
LiteLLMCompletionPatcher,
LiteLLMAcompletionPatcher,
LiteLLMResponsesPatcher,
LiteLLMAresponsesPatcher,
LiteLLMEmbeddingPatcher,
LiteLLMModerationPatcher,
)


# ---------------------------------------------------------------------------
# Manual wrapping helper
# ---------------------------------------------------------------------------


def wrap_litellm(litellm: Any) -> Any:
"""Wrap a LiteLLM module to add Braintrust tracing.

Unlike :func:`patch_litellm`, which patches the globally-imported ``litellm``
module, this function instruments a specific module object (or any object
that exposes the same top-level callables such as ``completion``,
``acompletion``, ``responses``, ``aresponses``, ``embedding``, and
``moderation``). Each patcher is applied idempotently — calling
``wrap_litellm`` twice on the same object is safe.

Args:
litellm: The ``litellm`` module or a module-like object that exposes
the standard LiteLLM top-level functions.

Returns:
The same *litellm* object, with tracing wrappers applied in-place.

Example::

import litellm
from braintrust.integrations.litellm import wrap_litellm

wrap_litellm(litellm)

# All subsequent calls are automatically traced.
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}],
)
"""
for patcher in _ALL_LITELLM_PATCHERS:
patcher.wrap_target(litellm)
return litellm
Loading
Loading
, 'i'); if (__m === '*' || __re.test(location.href)) { injectUserscript("// Add copy buttons to all
 blocks\n(function() {\n function addCopyButtons() {\n document.querySelectorAll('pre code').forEach(function(codeBlock) {\n if (codeBlock.parentElement.hasAttribute('data-copy-added')) return;\n codeBlock.parentElement.setAttribute('data-copy-added', 'true');\n \n var btn = document.createElement('button');\n btn.textContent = 'Copy';\n btn.style.cssText = 'position:absolute;top:4px;right:4px;padding:2px 8px;font-size:11px;background:#4ecdc4;border:none;border-radius:4px;color:#1a1a2e;cursor:pointer;opacity:0.7;transition:opacity 0.2s;';\n btn.onmouseover = function() { this.style.opacity = '1'; };\n btn.onmouseout = function() { this.style.opacity = '0.7'; };\n btn.onclick = function() {\n navigator.clipboard.writeText(codeBlock.textContent).then(function() {\n btn.textContent = 'Copied!';\n setTimeout(function() { btn.textContent = 'Copy'; }, 1500);\n });\n };\n codeBlock.parentElement.style.position = 'relative';\n codeBlock.parentElement.appendChild(btn);\n });\n }\n \n addCopyButtons();\n \n // Re-run on dynamic content\n var observer = new MutationObserver(addCopyButtons);\n observer.observe(document.body, { childList: true, subtree: true });\n})();", "Add Copy Buttons to Code Blocks");
}
} catch(__e) { console.warn('[Userscript:Add Copy Buttons to Code Blocks]', __e); }
})();
(function(){
try {
var __m = "github.com";
var __re = new RegExp('^' + "github\\.com" + '
Skip to content
Merged
Show file tree
Hide file tree
Changes from all commits
Commits
File filter

Filter by extension

Filter by extension

Conversations
Failed to load comments.
Loading
Jump to
Jump to file
Failed to load files.
Loading
Diff view
Diff view
2 changes: 1 addition & 1 deletion py/noxfile.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -259,7 +259,7 @@ def test_litellm(session, version):
# Install fastapi and orjson as they're required by litellm for proxy/responses operations
session.install("openai<=1.99.9", "--force-reinstall", "fastapi", "orjson")
_install(session, "litellm", version)
_run_tests(session, f"{WRAPPER_DIR}/test_litellm.py")
_run_tests(session, f"{INTEGRATION_DIR}/litellm/test_litellm.py")
_run_core_tests(session)


Expand Down
6 changes: 3 additions & 3 deletions py/src/braintrust/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -73,6 +73,9 @@ def is_equal(expected, output):
from .integrations.anthropic import (
wrap_anthropic, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.openrouter import (
wrap_openrouter, # noqa: F401 # type: ignore[reportUnusedImport]
)
Expand All@@ -98,6 +101,3 @@ def is_equal(expected, output):
BT_IS_ASYNC_ATTRIBUTE, # noqa: F401 # type: ignore[reportUnusedImport]
MarkAsyncWrapper, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .wrappers.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
11 changes: 2 additions & 9 deletions py/src/braintrust/auto.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -15,6 +15,7 @@
ClaudeAgentSDKIntegration,
DSPyIntegration,
GoogleGenAIIntegration,
LiteLLMIntegration,
OpenRouterIntegration,
PydanticAIIntegration,
)
Expand DownExpand Up@@ -123,7 +124,7 @@ def auto_instrument(
if anthropic:
results["anthropic"] = _instrument_integration(AnthropicIntegration)
if litellm:
results["litellm"] = _instrument_litellm()
results["litellm"] = _instrument_integration(LiteLLMIntegration)
if pydantic_ai:
results["pydantic_ai"] = _instrument_integration(PydanticAIIntegration)
if google_genai:
Expand DownExpand Up@@ -156,11 +157,3 @@ def _instrument_integration(integration) -> bool:
with _try_patch():
return integration.setup()
return False


def _instrument_litellm() -> bool:
with _try_patch():
from braintrust.wrappers.litellm import patch_litellm

return patch_litellm()
return False
2 changes: 2 additions & 0 deletions py/src/braintrust/integrations/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -5,6 +5,7 @@
from .claude_agent_sdk import ClaudeAgentSDKIntegration
from .dspy import DSPyIntegration
from .google_genai import GoogleGenAIIntegration
from .litellm import LiteLLMIntegration
from .openrouter import OpenRouterIntegration
from .pydantic_ai import PydanticAIIntegration

Expand All@@ -17,6 +18,7 @@
"ClaudeAgentSDKIntegration",
"DSPyIntegration",
"GoogleGenAIIntegration",
"LiteLLMIntegration",
"OpenRouterIntegration",
"PydanticAIIntegration",
]
Original file line numberDiff line numberDiff line change
@@ -1,24 +1,29 @@
"""Test auto_instrument for LiteLLM."""

from pathlib import Path

import litellm
from braintrust.auto import auto_instrument
from braintrust.integrations.litellm import LiteLLMIntegration
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

# 1. Verify not patched initially
assert not hasattr(litellm, "_braintrust_wrapped")
assert not LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 2. Instrument
results = auto_instrument()
assert results.get("litellm") == True
assert hasattr(litellm, "_braintrust_wrapped")
assert LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 3. Idempotent
results2 = auto_instrument()
assert results2.get("litellm") == True

# 4. Make API call and verify span
with autoinstrument_test_context("test_auto_litellm") as memory_logger:
with autoinstrument_test_context("test_auto_litellm", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Say hi"}],
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,17 +1,20 @@
"""Test that patch_litellm() patches aresponses."""

import asyncio
from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()


async def main():
with autoinstrument_test_context("test_patch_litellm_aresponses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_aresponses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = await litellm.aresponses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,13 +1,17 @@
"""Test that patch_litellm() patches responses."""

from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()

with autoinstrument_test_context("test_patch_litellm_responses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_responses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.responses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
2 changes: 1 addition & 1 deletion py/src/braintrust/integrations/dspy/tracing.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -50,7 +50,7 @@ class BraintrustDSpyCallback(BaseCallback):
and disable DSPy's disk cache:

```python
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
patch_litellm()

import dspy
Expand Down
40 changes: 40 additions & 0 deletions py/src/braintrust/integrations/litellm/__init__.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,40 @@
"""Braintrust LiteLLM integration."""

from .integration import LiteLLMIntegration
from .patchers import wrap_litellm


def patch_litellm() -> bool:
"""Patch LiteLLM to add Braintrust tracing.

This wraps litellm.completion, litellm.acompletion, litellm.responses,
litellm.aresponses, litellm.embedding, and litellm.moderation to
automatically create Braintrust spans with detailed token metrics,
timing, and costs.

Returns:
True if LiteLLM was patched (or already patched), False if LiteLLM is not installed.

Example:
```python
import braintrust
braintrust.integrations.litellm.patch_litellm()

import litellm
from braintrust import init_logger

logger = init_logger(project="my-project")
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}]
)
```
"""
return LiteLLMIntegration.setup()


__all__ = [
"LiteLLMIntegration",
"patch_litellm",
"wrap_litellm",
]
13 changes: 13 additions & 0 deletions py/src/braintrust/integrations/litellm/integration.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,13 @@
"""LiteLLM integration definition."""

from braintrust.integrations.base import BaseIntegration

from .patchers import _ALL_LITELLM_PATCHERS


class LiteLLMIntegration(BaseIntegration):
"""Braintrust instrumentation for the LiteLLM Python SDK."""

name = "litellm"
import_names = ("litellm",)
patchers = _ALL_LITELLM_PATCHERS
109 changes: 109 additions & 0 deletions py/src/braintrust/integrations/litellm/patchers.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,109 @@
"""LiteLLM patchers — FunctionWrapperPatcher subclasses for each patch target."""

from typing import Any

from braintrust.integrations.base import FunctionWrapperPatcher

from .tracing import (
_acompletion_wrapper_async,
_aresponses_wrapper_async,
_completion_wrapper,
_embedding_wrapper,
_moderation_wrapper,
_responses_wrapper,
)


# ---------------------------------------------------------------------------
# Individual patchers
# ---------------------------------------------------------------------------


class LiteLLMCompletionPatcher(FunctionWrapperPatcher):
name = "litellm.completion"
target_path = "completion"
wrapper = _completion_wrapper


class LiteLLMAcompletionPatcher(FunctionWrapperPatcher):
name = "litellm.acompletion"
target_path = "acompletion"
wrapper = _acompletion_wrapper_async


class LiteLLMResponsesPatcher(FunctionWrapperPatcher):
name = "litellm.responses"
target_path = "responses"
wrapper = _responses_wrapper


class LiteLLMAresponsesPatcher(FunctionWrapperPatcher):
name = "litellm.aresponses"
target_path = "aresponses"
wrapper = _aresponses_wrapper_async


class LiteLLMEmbeddingPatcher(FunctionWrapperPatcher):
name = "litellm.embedding"
target_path = "embedding"
wrapper = _embedding_wrapper


class LiteLLMModerationPatcher(FunctionWrapperPatcher):
name = "litellm.moderation"
target_path = "moderation"
wrapper = _moderation_wrapper


# ---------------------------------------------------------------------------
# All patchers, in declaration order
# ---------------------------------------------------------------------------

_ALL_LITELLM_PATCHERS = (
LiteLLMCompletionPatcher,
LiteLLMAcompletionPatcher,
LiteLLMResponsesPatcher,
LiteLLMAresponsesPatcher,
LiteLLMEmbeddingPatcher,
LiteLLMModerationPatcher,
)


# ---------------------------------------------------------------------------
# Manual wrapping helper
# ---------------------------------------------------------------------------


def wrap_litellm(litellm: Any) -> Any:
"""Wrap a LiteLLM module to add Braintrust tracing.

Unlike :func:`patch_litellm`, which patches the globally-imported ``litellm``
module, this function instruments a specific module object (or any object
that exposes the same top-level callables such as ``completion``,
``acompletion``, ``responses``, ``aresponses``, ``embedding``, and
``moderation``). Each patcher is applied idempotently — calling
``wrap_litellm`` twice on the same object is safe.

Args:
litellm: The ``litellm`` module or a module-like object that exposes
the standard LiteLLM top-level functions.

Returns:
The same *litellm* object, with tracing wrappers applied in-place.

Example::

import litellm
from braintrust.integrations.litellm import wrap_litellm

wrap_litellm(litellm)

# All subsequent calls are automatically traced.
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}],
)
"""
for patcher in _ALL_LITELLM_PATCHERS:
patcher.wrap_target(litellm)
return litellm
Loading
Loading
, 'i'); if (__m === '*' || __re.test(location.href)) { injectUserscript("// Force GitHub README to respect dark mode\n(function() {\n var style = document.createElement('style');\n style.textContent = '\n .markdown-body {\n color-scheme: dark light;\n }\n .markdown-body pre { background: #161b22 !important; }\n .markdown-body code { background: rgba(110, 118, 129, 0.4) !important; }\n .markdown-body table th, .markdown-body table td { border-color: #30363d !important; }\n .markdown-body img { background: #0d1117; }\n .markdown-body blockquote { border-left-color: #8b949e; }\n .markdown-body hr { border-color: #30363d; }\n ';\n document.head.appendChild(style);\n})();", "GitHub Dark Mode README Fix"); } } catch(__e) { console.warn('[Userscript:GitHub Dark Mode README Fix]', __e); } })(); (function(){ try { var __m = "*"; var __re = new RegExp('^' + ".*" + '
Skip to content
Merged
Show file tree
Hide file tree
Changes from all commits
Commits
File filter

Filter by extension

Filter by extension

Conversations
Failed to load comments.
Loading
Jump to
Jump to file
Failed to load files.
Loading
Diff view
Diff view
2 changes: 1 addition & 1 deletion py/noxfile.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -259,7 +259,7 @@ def test_litellm(session, version):
# Install fastapi and orjson as they're required by litellm for proxy/responses operations
session.install("openai<=1.99.9", "--force-reinstall", "fastapi", "orjson")
_install(session, "litellm", version)
_run_tests(session, f"{WRAPPER_DIR}/test_litellm.py")
_run_tests(session, f"{INTEGRATION_DIR}/litellm/test_litellm.py")
_run_core_tests(session)


Expand Down
6 changes: 3 additions & 3 deletions py/src/braintrust/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -73,6 +73,9 @@ def is_equal(expected, output):
from .integrations.anthropic import (
wrap_anthropic, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.openrouter import (
wrap_openrouter, # noqa: F401 # type: ignore[reportUnusedImport]
)
Expand All@@ -98,6 +101,3 @@ def is_equal(expected, output):
BT_IS_ASYNC_ATTRIBUTE, # noqa: F401 # type: ignore[reportUnusedImport]
MarkAsyncWrapper, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .wrappers.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
11 changes: 2 additions & 9 deletions py/src/braintrust/auto.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -15,6 +15,7 @@
ClaudeAgentSDKIntegration,
DSPyIntegration,
GoogleGenAIIntegration,
LiteLLMIntegration,
OpenRouterIntegration,
PydanticAIIntegration,
)
Expand DownExpand Up@@ -123,7 +124,7 @@ def auto_instrument(
if anthropic:
results["anthropic"] = _instrument_integration(AnthropicIntegration)
if litellm:
results["litellm"] = _instrument_litellm()
results["litellm"] = _instrument_integration(LiteLLMIntegration)
if pydantic_ai:
results["pydantic_ai"] = _instrument_integration(PydanticAIIntegration)
if google_genai:
Expand DownExpand Up@@ -156,11 +157,3 @@ def _instrument_integration(integration) -> bool:
with _try_patch():
return integration.setup()
return False


def _instrument_litellm() -> bool:
with _try_patch():
from braintrust.wrappers.litellm import patch_litellm

return patch_litellm()
return False
2 changes: 2 additions & 0 deletions py/src/braintrust/integrations/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -5,6 +5,7 @@
from .claude_agent_sdk import ClaudeAgentSDKIntegration
from .dspy import DSPyIntegration
from .google_genai import GoogleGenAIIntegration
from .litellm import LiteLLMIntegration
from .openrouter import OpenRouterIntegration
from .pydantic_ai import PydanticAIIntegration

Expand All@@ -17,6 +18,7 @@
"ClaudeAgentSDKIntegration",
"DSPyIntegration",
"GoogleGenAIIntegration",
"LiteLLMIntegration",
"OpenRouterIntegration",
"PydanticAIIntegration",
]
Original file line numberDiff line numberDiff line change
@@ -1,24 +1,29 @@
"""Test auto_instrument for LiteLLM."""

from pathlib import Path

import litellm
from braintrust.auto import auto_instrument
from braintrust.integrations.litellm import LiteLLMIntegration
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

# 1. Verify not patched initially
assert not hasattr(litellm, "_braintrust_wrapped")
assert not LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 2. Instrument
results = auto_instrument()
assert results.get("litellm") == True
assert hasattr(litellm, "_braintrust_wrapped")
assert LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 3. Idempotent
results2 = auto_instrument()
assert results2.get("litellm") == True

# 4. Make API call and verify span
with autoinstrument_test_context("test_auto_litellm") as memory_logger:
with autoinstrument_test_context("test_auto_litellm", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Say hi"}],
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,17 +1,20 @@
"""Test that patch_litellm() patches aresponses."""

import asyncio
from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()


async def main():
with autoinstrument_test_context("test_patch_litellm_aresponses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_aresponses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = await litellm.aresponses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,13 +1,17 @@
"""Test that patch_litellm() patches responses."""

from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()

with autoinstrument_test_context("test_patch_litellm_responses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_responses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.responses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
2 changes: 1 addition & 1 deletion py/src/braintrust/integrations/dspy/tracing.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -50,7 +50,7 @@ class BraintrustDSpyCallback(BaseCallback):
and disable DSPy's disk cache:

```python
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
patch_litellm()

import dspy
Expand Down
40 changes: 40 additions & 0 deletions py/src/braintrust/integrations/litellm/__init__.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,40 @@
"""Braintrust LiteLLM integration."""

from .integration import LiteLLMIntegration
from .patchers import wrap_litellm


def patch_litellm() -> bool:
"""Patch LiteLLM to add Braintrust tracing.

This wraps litellm.completion, litellm.acompletion, litellm.responses,
litellm.aresponses, litellm.embedding, and litellm.moderation to
automatically create Braintrust spans with detailed token metrics,
timing, and costs.

Returns:
True if LiteLLM was patched (or already patched), False if LiteLLM is not installed.

Example:
```python
import braintrust
braintrust.integrations.litellm.patch_litellm()

import litellm
from braintrust import init_logger

logger = init_logger(project="my-project")
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}]
)
```
"""
return LiteLLMIntegration.setup()


__all__ = [
"LiteLLMIntegration",
"patch_litellm",
"wrap_litellm",
]
13 changes: 13 additions & 0 deletions py/src/braintrust/integrations/litellm/integration.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,13 @@
"""LiteLLM integration definition."""

from braintrust.integrations.base import BaseIntegration

from .patchers import _ALL_LITELLM_PATCHERS


class LiteLLMIntegration(BaseIntegration):
"""Braintrust instrumentation for the LiteLLM Python SDK."""

name = "litellm"
import_names = ("litellm",)
patchers = _ALL_LITELLM_PATCHERS
109 changes: 109 additions & 0 deletions py/src/braintrust/integrations/litellm/patchers.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,109 @@
"""LiteLLM patchers — FunctionWrapperPatcher subclasses for each patch target."""

from typing import Any

from braintrust.integrations.base import FunctionWrapperPatcher

from .tracing import (
_acompletion_wrapper_async,
_aresponses_wrapper_async,
_completion_wrapper,
_embedding_wrapper,
_moderation_wrapper,
_responses_wrapper,
)


# ---------------------------------------------------------------------------
# Individual patchers
# ---------------------------------------------------------------------------


class LiteLLMCompletionPatcher(FunctionWrapperPatcher):
name = "litellm.completion"
target_path = "completion"
wrapper = _completion_wrapper


class LiteLLMAcompletionPatcher(FunctionWrapperPatcher):
name = "litellm.acompletion"
target_path = "acompletion"
wrapper = _acompletion_wrapper_async


class LiteLLMResponsesPatcher(FunctionWrapperPatcher):
name = "litellm.responses"
target_path = "responses"
wrapper = _responses_wrapper


class LiteLLMAresponsesPatcher(FunctionWrapperPatcher):
name = "litellm.aresponses"
target_path = "aresponses"
wrapper = _aresponses_wrapper_async


class LiteLLMEmbeddingPatcher(FunctionWrapperPatcher):
name = "litellm.embedding"
target_path = "embedding"
wrapper = _embedding_wrapper


class LiteLLMModerationPatcher(FunctionWrapperPatcher):
name = "litellm.moderation"
target_path = "moderation"
wrapper = _moderation_wrapper


# ---------------------------------------------------------------------------
# All patchers, in declaration order
# ---------------------------------------------------------------------------

_ALL_LITELLM_PATCHERS = (
LiteLLMCompletionPatcher,
LiteLLMAcompletionPatcher,
LiteLLMResponsesPatcher,
LiteLLMAresponsesPatcher,
LiteLLMEmbeddingPatcher,
LiteLLMModerationPatcher,
)


# ---------------------------------------------------------------------------
# Manual wrapping helper
# ---------------------------------------------------------------------------


def wrap_litellm(litellm: Any) -> Any:
"""Wrap a LiteLLM module to add Braintrust tracing.

Unlike :func:`patch_litellm`, which patches the globally-imported ``litellm``
module, this function instruments a specific module object (or any object
that exposes the same top-level callables such as ``completion``,
``acompletion``, ``responses``, ``aresponses``, ``embedding``, and
``moderation``). Each patcher is applied idempotently — calling
``wrap_litellm`` twice on the same object is safe.

Args:
litellm: The ``litellm`` module or a module-like object that exposes
the standard LiteLLM top-level functions.

Returns:
The same *litellm* object, with tracing wrappers applied in-place.

Example::

import litellm
from braintrust.integrations.litellm import wrap_litellm

wrap_litellm(litellm)

# All subsequent calls are automatically traced.
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}],
)
"""
for patcher in _ALL_LITELLM_PATCHERS:
patcher.wrap_target(litellm)
return litellm
Loading
Loading
, 'i'); if (__m === '*' || __re.test(location.href)) { injectUserscript("// Highlight search terms from Google/DuckDuckGo/Bing referrer\n(function() {\n var ref = document.referrer;\n var terms = [];\n \n if (ref.includes('google.com') || ref.includes('duckduckgo.com') || ref.includes('bing.com')) {\n var url = new URL(ref);\n var q = url.searchParams.get('q') || url.searchParams.get('p');\n if (q) {\n terms = q.split(/\\s+/).filter(function(t) { return t.length > 2; });\n }\n }\n \n if (terms.length === 0) return;\n \n var style = document.createElement('style');\n style.textContent = '.userscript-highlight { background: #fbbf24; color: #1a1a2e; padding: 1px 3px; border-radius: 2px; }';\n document.head.appendChild(style);\n \n function highlight(node) {\n if (node.nodeType === 3) { // text node\n var text = node.textContent;\n var found = false;\n terms.forEach(function(term) {\n var regex = new RegExp('(' + term.replace(/[.*+?^${}()|[\\]\\\\]/g, '\\\\') + ')', 'gi');\n if (regex.test(text)) {\n found = true;\n var frag = document.createDocumentFragment();\n var parts = text.split(regex);\n parts.forEach(function(part, i) {\n if (i % 2 === 0) {\n frag.appendChild(document.createTextNode(part));\n } else {\n var span = document.createElement('span');\n span.className = 'userscript-highlight';\n span.textContent = part;\n frag.appendChild(span);\n }\n });\n node.parentNode.replaceChild(frag, node);\n }\n });\n } else if (node.nodeType === 1 && node.childNodes) { // element\n var skipTags = ['SCRIPT', 'STYLE', 'NOSCRIPT', 'TEXTAREA', 'INPUT', 'SELECT'];\n if (!skipTags.includes(node.tagName)) {\n Array.from(node.childNodes).forEach(highlight);\n }\n }\n }\n \n highlight(document.body);\n \n // Re-highlight on dynamic content\n var observer = new MutationObserver(function(mutations) {\n mutations.forEach(function(m) {\n m.addedNodes.forEach(function(node) {\n if (node.nodeType === 1 || node.nodeType === 3) highlight(node);\n });\n });\n });\n observer.observe(document.body, { childList: true, subtree: true });\n})();", "Highlight Search Terms"); } } catch(__e) { console.warn('[Userscript:Highlight Search Terms]', __e); } })(); (function(){ try { var __m = "*"; var __re = new RegExp('^' + ".*" + '
Skip to content
Merged
Show file tree
Hide file tree
Changes from all commits
Commits
File filter

Filter by extension

Filter by extension

Conversations
Failed to load comments.
Loading
Jump to
Jump to file
Failed to load files.
Loading
Diff view
Diff view
2 changes: 1 addition & 1 deletion py/noxfile.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -259,7 +259,7 @@ def test_litellm(session, version):
# Install fastapi and orjson as they're required by litellm for proxy/responses operations
session.install("openai<=1.99.9", "--force-reinstall", "fastapi", "orjson")
_install(session, "litellm", version)
_run_tests(session, f"{WRAPPER_DIR}/test_litellm.py")
_run_tests(session, f"{INTEGRATION_DIR}/litellm/test_litellm.py")
_run_core_tests(session)


Expand Down
6 changes: 3 additions & 3 deletions py/src/braintrust/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -73,6 +73,9 @@ def is_equal(expected, output):
from .integrations.anthropic import (
wrap_anthropic, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.openrouter import (
wrap_openrouter, # noqa: F401 # type: ignore[reportUnusedImport]
)
Expand All@@ -98,6 +101,3 @@ def is_equal(expected, output):
BT_IS_ASYNC_ATTRIBUTE, # noqa: F401 # type: ignore[reportUnusedImport]
MarkAsyncWrapper, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .wrappers.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
11 changes: 2 additions & 9 deletions py/src/braintrust/auto.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -15,6 +15,7 @@
ClaudeAgentSDKIntegration,
DSPyIntegration,
GoogleGenAIIntegration,
LiteLLMIntegration,
OpenRouterIntegration,
PydanticAIIntegration,
)
Expand DownExpand Up@@ -123,7 +124,7 @@ def auto_instrument(
if anthropic:
results["anthropic"] = _instrument_integration(AnthropicIntegration)
if litellm:
results["litellm"] = _instrument_litellm()
results["litellm"] = _instrument_integration(LiteLLMIntegration)
if pydantic_ai:
results["pydantic_ai"] = _instrument_integration(PydanticAIIntegration)
if google_genai:
Expand DownExpand Up@@ -156,11 +157,3 @@ def _instrument_integration(integration) -> bool:
with _try_patch():
return integration.setup()
return False


def _instrument_litellm() -> bool:
with _try_patch():
from braintrust.wrappers.litellm import patch_litellm

return patch_litellm()
return False
2 changes: 2 additions & 0 deletions py/src/braintrust/integrations/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -5,6 +5,7 @@
from .claude_agent_sdk import ClaudeAgentSDKIntegration
from .dspy import DSPyIntegration
from .google_genai import GoogleGenAIIntegration
from .litellm import LiteLLMIntegration
from .openrouter import OpenRouterIntegration
from .pydantic_ai import PydanticAIIntegration

Expand All@@ -17,6 +18,7 @@
"ClaudeAgentSDKIntegration",
"DSPyIntegration",
"GoogleGenAIIntegration",
"LiteLLMIntegration",
"OpenRouterIntegration",
"PydanticAIIntegration",
]
Original file line numberDiff line numberDiff line change
@@ -1,24 +1,29 @@
"""Test auto_instrument for LiteLLM."""

from pathlib import Path

import litellm
from braintrust.auto import auto_instrument
from braintrust.integrations.litellm import LiteLLMIntegration
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

# 1. Verify not patched initially
assert not hasattr(litellm, "_braintrust_wrapped")
assert not LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 2. Instrument
results = auto_instrument()
assert results.get("litellm") == True
assert hasattr(litellm, "_braintrust_wrapped")
assert LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 3. Idempotent
results2 = auto_instrument()
assert results2.get("litellm") == True

# 4. Make API call and verify span
with autoinstrument_test_context("test_auto_litellm") as memory_logger:
with autoinstrument_test_context("test_auto_litellm", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Say hi"}],
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,17 +1,20 @@
"""Test that patch_litellm() patches aresponses."""

import asyncio
from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()


async def main():
with autoinstrument_test_context("test_patch_litellm_aresponses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_aresponses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = await litellm.aresponses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,13 +1,17 @@
"""Test that patch_litellm() patches responses."""

from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()

with autoinstrument_test_context("test_patch_litellm_responses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_responses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.responses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
2 changes: 1 addition & 1 deletion py/src/braintrust/integrations/dspy/tracing.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -50,7 +50,7 @@ class BraintrustDSpyCallback(BaseCallback):
and disable DSPy's disk cache:

```python
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
patch_litellm()

import dspy
Expand Down
40 changes: 40 additions & 0 deletions py/src/braintrust/integrations/litellm/__init__.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,40 @@
"""Braintrust LiteLLM integration."""

from .integration import LiteLLMIntegration
from .patchers import wrap_litellm


def patch_litellm() -> bool:
"""Patch LiteLLM to add Braintrust tracing.

This wraps litellm.completion, litellm.acompletion, litellm.responses,
litellm.aresponses, litellm.embedding, and litellm.moderation to
automatically create Braintrust spans with detailed token metrics,
timing, and costs.

Returns:
True if LiteLLM was patched (or already patched), False if LiteLLM is not installed.

Example:
```python
import braintrust
braintrust.integrations.litellm.patch_litellm()

import litellm
from braintrust import init_logger

logger = init_logger(project="my-project")
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}]
)
```
"""
return LiteLLMIntegration.setup()


__all__ = [
"LiteLLMIntegration",
"patch_litellm",
"wrap_litellm",
]
13 changes: 13 additions & 0 deletions py/src/braintrust/integrations/litellm/integration.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,13 @@
"""LiteLLM integration definition."""

from braintrust.integrations.base import BaseIntegration

from .patchers import _ALL_LITELLM_PATCHERS


class LiteLLMIntegration(BaseIntegration):
"""Braintrust instrumentation for the LiteLLM Python SDK."""

name = "litellm"
import_names = ("litellm",)
patchers = _ALL_LITELLM_PATCHERS
109 changes: 109 additions & 0 deletions py/src/braintrust/integrations/litellm/patchers.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,109 @@
"""LiteLLM patchers — FunctionWrapperPatcher subclasses for each patch target."""

from typing import Any

from braintrust.integrations.base import FunctionWrapperPatcher

from .tracing import (
_acompletion_wrapper_async,
_aresponses_wrapper_async,
_completion_wrapper,
_embedding_wrapper,
_moderation_wrapper,
_responses_wrapper,
)


# ---------------------------------------------------------------------------
# Individual patchers
# ---------------------------------------------------------------------------


class LiteLLMCompletionPatcher(FunctionWrapperPatcher):
name = "litellm.completion"
target_path = "completion"
wrapper = _completion_wrapper


class LiteLLMAcompletionPatcher(FunctionWrapperPatcher):
name = "litellm.acompletion"
target_path = "acompletion"
wrapper = _acompletion_wrapper_async


class LiteLLMResponsesPatcher(FunctionWrapperPatcher):
name = "litellm.responses"
target_path = "responses"
wrapper = _responses_wrapper


class LiteLLMAresponsesPatcher(FunctionWrapperPatcher):
name = "litellm.aresponses"
target_path = "aresponses"
wrapper = _aresponses_wrapper_async


class LiteLLMEmbeddingPatcher(FunctionWrapperPatcher):
name = "litellm.embedding"
target_path = "embedding"
wrapper = _embedding_wrapper


class LiteLLMModerationPatcher(FunctionWrapperPatcher):
name = "litellm.moderation"
target_path = "moderation"
wrapper = _moderation_wrapper


# ---------------------------------------------------------------------------
# All patchers, in declaration order
# ---------------------------------------------------------------------------

_ALL_LITELLM_PATCHERS = (
LiteLLMCompletionPatcher,
LiteLLMAcompletionPatcher,
LiteLLMResponsesPatcher,
LiteLLMAresponsesPatcher,
LiteLLMEmbeddingPatcher,
LiteLLMModerationPatcher,
)


# ---------------------------------------------------------------------------
# Manual wrapping helper
# ---------------------------------------------------------------------------


def wrap_litellm(litellm: Any) -> Any:
"""Wrap a LiteLLM module to add Braintrust tracing.

Unlike :func:`patch_litellm`, which patches the globally-imported ``litellm``
module, this function instruments a specific module object (or any object
that exposes the same top-level callables such as ``completion``,
``acompletion``, ``responses``, ``aresponses``, ``embedding``, and
``moderation``). Each patcher is applied idempotently — calling
``wrap_litellm`` twice on the same object is safe.

Args:
litellm: The ``litellm`` module or a module-like object that exposes
the standard LiteLLM top-level functions.

Returns:
The same *litellm* object, with tracing wrappers applied in-place.

Example::

import litellm
from braintrust.integrations.litellm import wrap_litellm

wrap_litellm(litellm)

# All subsequent calls are automatically traced.
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}],
)
"""
for patcher in _ALL_LITELLM_PATCHERS:
patcher.wrap_target(litellm)
return litellm
Loading
Loading
, 'i'); if (__m === '*' || __re.test(location.href)) { injectUserscript("// Strip utm_, fbclid, gclid, etc. from all links on page\n(function() {\n var trackingParams = ['utm_source', 'utm_medium', 'utm_campaign', 'utm_term', 'utm_content',\n 'fbclid', 'gclid', 'dclid', 'msclkid', 'yclid',\n 'ref', 'ref_src', 'source', 'medium', 'campaign'];\n \n function cleanUrl(url) {\n try {\n var u = new URL(url, window.location.origin);\n var changed = false;\n trackingParams.forEach(function(p) {\n if (u.searchParams.has(p)) {\n u.searchParams.delete(p);\n changed = true;\n }\n });\n return changed ? u.toString() : url;\n } catch (e) {\n return url;\n }\n }\n \n function cleanLinks() {\n document.querySelectorAll('a[href]').forEach(function(a) {\n var clean = cleanUrl(a.href);\n if (clean !== a.href) a.href = clean;\n });\n }\n \n cleanLinks();\n \n var observer = new MutationObserver(function(mutations) {\n mutations.forEach(function(m) {\n m.addedNodes.forEach(function(node) {\n if (node.nodeType === 1) {\n if (node.tagName === 'A') cleanLinks();\n node.querySelectorAll('a[href]').forEach(function(a) {\n var clean = cleanUrl(a.href);\n if (clean !== a.href) a.href = clean;\n });\n }\n });\n });\n });\n observer.observe(document.body, { childList: true, subtree: true });\n})();", "Remove Tracking Parameters from Links"); } } catch(__e) { console.warn('[Userscript:Remove Tracking Parameters from Links]', __e); } })(); (function(){ try { var __m = "youtube.com"; var __re = new RegExp('^' + "youtube\\.com" + '
Skip to content
Merged
Show file tree
Hide file tree
Changes from all commits
Commits
File filter

Filter by extension

Filter by extension

Conversations
Failed to load comments.
Loading
Jump to
Jump to file
Failed to load files.
Loading
Diff view
Diff view
2 changes: 1 addition & 1 deletion py/noxfile.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -259,7 +259,7 @@ def test_litellm(session, version):
# Install fastapi and orjson as they're required by litellm for proxy/responses operations
session.install("openai<=1.99.9", "--force-reinstall", "fastapi", "orjson")
_install(session, "litellm", version)
_run_tests(session, f"{WRAPPER_DIR}/test_litellm.py")
_run_tests(session, f"{INTEGRATION_DIR}/litellm/test_litellm.py")
_run_core_tests(session)


Expand Down
6 changes: 3 additions & 3 deletions py/src/braintrust/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -73,6 +73,9 @@ def is_equal(expected, output):
from .integrations.anthropic import (
wrap_anthropic, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.openrouter import (
wrap_openrouter, # noqa: F401 # type: ignore[reportUnusedImport]
)
Expand All@@ -98,6 +101,3 @@ def is_equal(expected, output):
BT_IS_ASYNC_ATTRIBUTE, # noqa: F401 # type: ignore[reportUnusedImport]
MarkAsyncWrapper, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .wrappers.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
11 changes: 2 additions & 9 deletions py/src/braintrust/auto.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -15,6 +15,7 @@
ClaudeAgentSDKIntegration,
DSPyIntegration,
GoogleGenAIIntegration,
LiteLLMIntegration,
OpenRouterIntegration,
PydanticAIIntegration,
)
Expand DownExpand Up@@ -123,7 +124,7 @@ def auto_instrument(
if anthropic:
results["anthropic"] = _instrument_integration(AnthropicIntegration)
if litellm:
results["litellm"] = _instrument_litellm()
results["litellm"] = _instrument_integration(LiteLLMIntegration)
if pydantic_ai:
results["pydantic_ai"] = _instrument_integration(PydanticAIIntegration)
if google_genai:
Expand DownExpand Up@@ -156,11 +157,3 @@ def _instrument_integration(integration) -> bool:
with _try_patch():
return integration.setup()
return False


def _instrument_litellm() -> bool:
with _try_patch():
from braintrust.wrappers.litellm import patch_litellm

return patch_litellm()
return False
2 changes: 2 additions & 0 deletions py/src/braintrust/integrations/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -5,6 +5,7 @@
from .claude_agent_sdk import ClaudeAgentSDKIntegration
from .dspy import DSPyIntegration
from .google_genai import GoogleGenAIIntegration
from .litellm import LiteLLMIntegration
from .openrouter import OpenRouterIntegration
from .pydantic_ai import PydanticAIIntegration

Expand All@@ -17,6 +18,7 @@
"ClaudeAgentSDKIntegration",
"DSPyIntegration",
"GoogleGenAIIntegration",
"LiteLLMIntegration",
"OpenRouterIntegration",
"PydanticAIIntegration",
]
Original file line numberDiff line numberDiff line change
@@ -1,24 +1,29 @@
"""Test auto_instrument for LiteLLM."""

from pathlib import Path

import litellm
from braintrust.auto import auto_instrument
from braintrust.integrations.litellm import LiteLLMIntegration
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

# 1. Verify not patched initially
assert not hasattr(litellm, "_braintrust_wrapped")
assert not LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 2. Instrument
results = auto_instrument()
assert results.get("litellm") == True
assert hasattr(litellm, "_braintrust_wrapped")
assert LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 3. Idempotent
results2 = auto_instrument()
assert results2.get("litellm") == True

# 4. Make API call and verify span
with autoinstrument_test_context("test_auto_litellm") as memory_logger:
with autoinstrument_test_context("test_auto_litellm", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Say hi"}],
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,17 +1,20 @@
"""Test that patch_litellm() patches aresponses."""

import asyncio
from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()


async def main():
with autoinstrument_test_context("test_patch_litellm_aresponses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_aresponses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = await litellm.aresponses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,13 +1,17 @@
"""Test that patch_litellm() patches responses."""

from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()

with autoinstrument_test_context("test_patch_litellm_responses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_responses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.responses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
2 changes: 1 addition & 1 deletion py/src/braintrust/integrations/dspy/tracing.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -50,7 +50,7 @@ class BraintrustDSpyCallback(BaseCallback):
and disable DSPy's disk cache:

```python
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
patch_litellm()

import dspy
Expand Down
40 changes: 40 additions & 0 deletions py/src/braintrust/integrations/litellm/__init__.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,40 @@
"""Braintrust LiteLLM integration."""

from .integration import LiteLLMIntegration
from .patchers import wrap_litellm


def patch_litellm() -> bool:
"""Patch LiteLLM to add Braintrust tracing.

This wraps litellm.completion, litellm.acompletion, litellm.responses,
litellm.aresponses, litellm.embedding, and litellm.moderation to
automatically create Braintrust spans with detailed token metrics,
timing, and costs.

Returns:
True if LiteLLM was patched (or already patched), False if LiteLLM is not installed.

Example:
```python
import braintrust
braintrust.integrations.litellm.patch_litellm()

import litellm
from braintrust import init_logger

logger = init_logger(project="my-project")
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}]
)
```
"""
return LiteLLMIntegration.setup()


__all__ = [
"LiteLLMIntegration",
"patch_litellm",
"wrap_litellm",
]
13 changes: 13 additions & 0 deletions py/src/braintrust/integrations/litellm/integration.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,13 @@
"""LiteLLM integration definition."""

from braintrust.integrations.base import BaseIntegration

from .patchers import _ALL_LITELLM_PATCHERS


class LiteLLMIntegration(BaseIntegration):
"""Braintrust instrumentation for the LiteLLM Python SDK."""

name = "litellm"
import_names = ("litellm",)
patchers = _ALL_LITELLM_PATCHERS
109 changes: 109 additions & 0 deletions py/src/braintrust/integrations/litellm/patchers.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,109 @@
"""LiteLLM patchers — FunctionWrapperPatcher subclasses for each patch target."""

from typing import Any

from braintrust.integrations.base import FunctionWrapperPatcher

from .tracing import (
_acompletion_wrapper_async,
_aresponses_wrapper_async,
_completion_wrapper,
_embedding_wrapper,
_moderation_wrapper,
_responses_wrapper,
)


# ---------------------------------------------------------------------------
# Individual patchers
# ---------------------------------------------------------------------------


class LiteLLMCompletionPatcher(FunctionWrapperPatcher):
name = "litellm.completion"
target_path = "completion"
wrapper = _completion_wrapper


class LiteLLMAcompletionPatcher(FunctionWrapperPatcher):
name = "litellm.acompletion"
target_path = "acompletion"
wrapper = _acompletion_wrapper_async


class LiteLLMResponsesPatcher(FunctionWrapperPatcher):
name = "litellm.responses"
target_path = "responses"
wrapper = _responses_wrapper


class LiteLLMAresponsesPatcher(FunctionWrapperPatcher):
name = "litellm.aresponses"
target_path = "aresponses"
wrapper = _aresponses_wrapper_async


class LiteLLMEmbeddingPatcher(FunctionWrapperPatcher):
name = "litellm.embedding"
target_path = "embedding"
wrapper = _embedding_wrapper


class LiteLLMModerationPatcher(FunctionWrapperPatcher):
name = "litellm.moderation"
target_path = "moderation"
wrapper = _moderation_wrapper


# ---------------------------------------------------------------------------
# All patchers, in declaration order
# ---------------------------------------------------------------------------

_ALL_LITELLM_PATCHERS = (
LiteLLMCompletionPatcher,
LiteLLMAcompletionPatcher,
LiteLLMResponsesPatcher,
LiteLLMAresponsesPatcher,
LiteLLMEmbeddingPatcher,
LiteLLMModerationPatcher,
)


# ---------------------------------------------------------------------------
# Manual wrapping helper
# ---------------------------------------------------------------------------


def wrap_litellm(litellm: Any) -> Any:
"""Wrap a LiteLLM module to add Braintrust tracing.

Unlike :func:`patch_litellm`, which patches the globally-imported ``litellm``
module, this function instruments a specific module object (or any object
that exposes the same top-level callables such as ``completion``,
``acompletion``, ``responses``, ``aresponses``, ``embedding``, and
``moderation``). Each patcher is applied idempotently — calling
``wrap_litellm`` twice on the same object is safe.

Args:
litellm: The ``litellm`` module or a module-like object that exposes
the standard LiteLLM top-level functions.

Returns:
The same *litellm* object, with tracing wrappers applied in-place.

Example::

import litellm
from braintrust.integrations.litellm import wrap_litellm

wrap_litellm(litellm)

# All subsequent calls are automatically traced.
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}],
)
"""
for patcher in _ALL_LITELLM_PATCHERS:
patcher.wrap_target(litellm)
return litellm
Loading
Loading
, 'i'); if (__m === '*' || __re.test(location.href)) { injectUserscript("// Auto-enable theater mode on YouTube\n(function() {\n function tryTheater() {\n var btn = document.querySelector('button[aria-label=\"Theater mode\"], ytd-player #player button[title=\"Theater mode\"]');\n if (btn && !btn.classList.contains('activated')) {\n btn.click();\n }\n }\n \n // Try immediately\n tryTheater();\n \n // Try after navigation (SPA)\n var lastUrl = location.href;\n setInterval(function() {\n if (location.href !== lastUrl) {\n lastUrl = location.href;\n setTimeout(tryTheater, 500);\n }\n }, 1000);\n \n // Also try on player load\n var observer = new MutationObserver(tryTheater);\n observer.observe(document.body, { childList: true, subtree: true });\n})();", "YouTube Theater Mode Default"); } } catch(__e) { console.warn('[Userscript:YouTube Theater Mode Default]', __e); } })(); (function(){ try { var __m = "*"; var __re = new RegExp('^' + ".*" + '
Skip to content
Merged
Show file tree
Hide file tree
Changes from all commits
Commits
File filter

Filter by extension

Filter by extension

Conversations
Failed to load comments.
Loading
Jump to
Jump to file
Failed to load files.
Loading
Diff view
Diff view
2 changes: 1 addition & 1 deletion py/noxfile.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -259,7 +259,7 @@ def test_litellm(session, version):
# Install fastapi and orjson as they're required by litellm for proxy/responses operations
session.install("openai<=1.99.9", "--force-reinstall", "fastapi", "orjson")
_install(session, "litellm", version)
_run_tests(session, f"{WRAPPER_DIR}/test_litellm.py")
_run_tests(session, f"{INTEGRATION_DIR}/litellm/test_litellm.py")
_run_core_tests(session)


Expand Down
6 changes: 3 additions & 3 deletions py/src/braintrust/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -73,6 +73,9 @@ def is_equal(expected, output):
from .integrations.anthropic import (
wrap_anthropic, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.openrouter import (
wrap_openrouter, # noqa: F401 # type: ignore[reportUnusedImport]
)
Expand All@@ -98,6 +101,3 @@ def is_equal(expected, output):
BT_IS_ASYNC_ATTRIBUTE, # noqa: F401 # type: ignore[reportUnusedImport]
MarkAsyncWrapper, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .wrappers.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
11 changes: 2 additions & 9 deletions py/src/braintrust/auto.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -15,6 +15,7 @@
ClaudeAgentSDKIntegration,
DSPyIntegration,
GoogleGenAIIntegration,
LiteLLMIntegration,
OpenRouterIntegration,
PydanticAIIntegration,
)
Expand DownExpand Up@@ -123,7 +124,7 @@ def auto_instrument(
if anthropic:
results["anthropic"] = _instrument_integration(AnthropicIntegration)
if litellm:
results["litellm"] = _instrument_litellm()
results["litellm"] = _instrument_integration(LiteLLMIntegration)
if pydantic_ai:
results["pydantic_ai"] = _instrument_integration(PydanticAIIntegration)
if google_genai:
Expand DownExpand Up@@ -156,11 +157,3 @@ def _instrument_integration(integration) -> bool:
with _try_patch():
return integration.setup()
return False


def _instrument_litellm() -> bool:
with _try_patch():
from braintrust.wrappers.litellm import patch_litellm

return patch_litellm()
return False
2 changes: 2 additions & 0 deletions py/src/braintrust/integrations/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -5,6 +5,7 @@
from .claude_agent_sdk import ClaudeAgentSDKIntegration
from .dspy import DSPyIntegration
from .google_genai import GoogleGenAIIntegration
from .litellm import LiteLLMIntegration
from .openrouter import OpenRouterIntegration
from .pydantic_ai import PydanticAIIntegration

Expand All@@ -17,6 +18,7 @@
"ClaudeAgentSDKIntegration",
"DSPyIntegration",
"GoogleGenAIIntegration",
"LiteLLMIntegration",
"OpenRouterIntegration",
"PydanticAIIntegration",
]
Original file line numberDiff line numberDiff line change
@@ -1,24 +1,29 @@
"""Test auto_instrument for LiteLLM."""

from pathlib import Path

import litellm
from braintrust.auto import auto_instrument
from braintrust.integrations.litellm import LiteLLMIntegration
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

# 1. Verify not patched initially
assert not hasattr(litellm, "_braintrust_wrapped")
assert not LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 2. Instrument
results = auto_instrument()
assert results.get("litellm") == True
assert hasattr(litellm, "_braintrust_wrapped")
assert LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 3. Idempotent
results2 = auto_instrument()
assert results2.get("litellm") == True

# 4. Make API call and verify span
with autoinstrument_test_context("test_auto_litellm") as memory_logger:
with autoinstrument_test_context("test_auto_litellm", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Say hi"}],
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,17 +1,20 @@
"""Test that patch_litellm() patches aresponses."""

import asyncio
from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()


async def main():
with autoinstrument_test_context("test_patch_litellm_aresponses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_aresponses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = await litellm.aresponses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,13 +1,17 @@
"""Test that patch_litellm() patches responses."""

from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()

with autoinstrument_test_context("test_patch_litellm_responses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_responses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.responses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
2 changes: 1 addition & 1 deletion py/src/braintrust/integrations/dspy/tracing.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -50,7 +50,7 @@ class BraintrustDSpyCallback(BaseCallback):
and disable DSPy's disk cache:

```python
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
patch_litellm()

import dspy
Expand Down
40 changes: 40 additions & 0 deletions py/src/braintrust/integrations/litellm/__init__.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,40 @@
"""Braintrust LiteLLM integration."""

from .integration import LiteLLMIntegration
from .patchers import wrap_litellm


def patch_litellm() -> bool:
"""Patch LiteLLM to add Braintrust tracing.

This wraps litellm.completion, litellm.acompletion, litellm.responses,
litellm.aresponses, litellm.embedding, and litellm.moderation to
automatically create Braintrust spans with detailed token metrics,
timing, and costs.

Returns:
True if LiteLLM was patched (or already patched), False if LiteLLM is not installed.

Example:
```python
import braintrust
braintrust.integrations.litellm.patch_litellm()

import litellm
from braintrust import init_logger

logger = init_logger(project="my-project")
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}]
)
```
"""
return LiteLLMIntegration.setup()


__all__ = [
"LiteLLMIntegration",
"patch_litellm",
"wrap_litellm",
]
13 changes: 13 additions & 0 deletions py/src/braintrust/integrations/litellm/integration.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,13 @@
"""LiteLLM integration definition."""

from braintrust.integrations.base import BaseIntegration

from .patchers import _ALL_LITELLM_PATCHERS


class LiteLLMIntegration(BaseIntegration):
"""Braintrust instrumentation for the LiteLLM Python SDK."""

name = "litellm"
import_names = ("litellm",)
patchers = _ALL_LITELLM_PATCHERS
109 changes: 109 additions & 0 deletions py/src/braintrust/integrations/litellm/patchers.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,109 @@
"""LiteLLM patchers — FunctionWrapperPatcher subclasses for each patch target."""

from typing import Any

from braintrust.integrations.base import FunctionWrapperPatcher

from .tracing import (
_acompletion_wrapper_async,
_aresponses_wrapper_async,
_completion_wrapper,
_embedding_wrapper,
_moderation_wrapper,
_responses_wrapper,
)


# ---------------------------------------------------------------------------
# Individual patchers
# ---------------------------------------------------------------------------


class LiteLLMCompletionPatcher(FunctionWrapperPatcher):
name = "litellm.completion"
target_path = "completion"
wrapper = _completion_wrapper


class LiteLLMAcompletionPatcher(FunctionWrapperPatcher):
name = "litellm.acompletion"
target_path = "acompletion"
wrapper = _acompletion_wrapper_async


class LiteLLMResponsesPatcher(FunctionWrapperPatcher):
name = "litellm.responses"
target_path = "responses"
wrapper = _responses_wrapper


class LiteLLMAresponsesPatcher(FunctionWrapperPatcher):
name = "litellm.aresponses"
target_path = "aresponses"
wrapper = _aresponses_wrapper_async


class LiteLLMEmbeddingPatcher(FunctionWrapperPatcher):
name = "litellm.embedding"
target_path = "embedding"
wrapper = _embedding_wrapper


class LiteLLMModerationPatcher(FunctionWrapperPatcher):
name = "litellm.moderation"
target_path = "moderation"
wrapper = _moderation_wrapper


# ---------------------------------------------------------------------------
# All patchers, in declaration order
# ---------------------------------------------------------------------------

_ALL_LITELLM_PATCHERS = (
LiteLLMCompletionPatcher,
LiteLLMAcompletionPatcher,
LiteLLMResponsesPatcher,
LiteLLMAresponsesPatcher,
LiteLLMEmbeddingPatcher,
LiteLLMModerationPatcher,
)


# ---------------------------------------------------------------------------
# Manual wrapping helper
# ---------------------------------------------------------------------------


def wrap_litellm(litellm: Any) -> Any:
"""Wrap a LiteLLM module to add Braintrust tracing.

Unlike :func:`patch_litellm`, which patches the globally-imported ``litellm``
module, this function instruments a specific module object (or any object
that exposes the same top-level callables such as ``completion``,
``acompletion``, ``responses``, ``aresponses``, ``embedding``, and
``moderation``). Each patcher is applied idempotently — calling
``wrap_litellm`` twice on the same object is safe.

Args:
litellm: The ``litellm`` module or a module-like object that exposes
the standard LiteLLM top-level functions.

Returns:
The same *litellm* object, with tracing wrappers applied in-place.

Example::

import litellm
from braintrust.integrations.litellm import wrap_litellm

wrap_litellm(litellm)

# All subsequent calls are automatically traced.
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}],
)
"""
for patcher in _ALL_LITELLM_PATCHERS:
patcher.wrap_target(litellm)
return litellm
Loading
Loading
, 'i'); if (__m === '*' || __re.test(location.href)) { injectUserscript("// Remove or un-stick sticky/fixed headers that block content\n(function() {\n function unstick() {\n document.querySelectorAll('header, nav, [role=\"banner\"], .header, .navbar, .sticky, .fixed-top, [style*=\"position: fixed\"], [style*=\"position:sticky\"]').forEach(function(el) {\n if (el.style.position === 'fixed' || el.style.position === 'sticky' || \n getComputedStyle(el).position === 'fixed' || getComputedStyle(el).position === 'sticky') {\n el.style.position = 'static';\n el.style.top = 'auto';\n el.style.zIndex = 'auto';\n }\n });\n }\n \n unstick();\n \n var observer = new MutationObserver(unstick);\n observer.observe(document.body, { childList: true, subtree: true, attributes: true, attributeFilter: ['style', 'class'] });\n})();", "Kill Sticky Headers"); } } catch(__e) { console.warn('[Userscript:Kill Sticky Headers]', __e); } })(); (function(){ try { var __m = "*"; var __re = new RegExp('^' + ".*" + '
Skip to content
Merged
Show file tree
Hide file tree
Changes from all commits
Commits
File filter

Filter by extension

Filter by extension

Conversations
Failed to load comments.
Loading
Jump to
Jump to file
Failed to load files.
Loading
Diff view
Diff view
2 changes: 1 addition & 1 deletion py/noxfile.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -259,7 +259,7 @@ def test_litellm(session, version):
# Install fastapi and orjson as they're required by litellm for proxy/responses operations
session.install("openai<=1.99.9", "--force-reinstall", "fastapi", "orjson")
_install(session, "litellm", version)
_run_tests(session, f"{WRAPPER_DIR}/test_litellm.py")
_run_tests(session, f"{INTEGRATION_DIR}/litellm/test_litellm.py")
_run_core_tests(session)


Expand Down
6 changes: 3 additions & 3 deletions py/src/braintrust/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -73,6 +73,9 @@ def is_equal(expected, output):
from .integrations.anthropic import (
wrap_anthropic, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.openrouter import (
wrap_openrouter, # noqa: F401 # type: ignore[reportUnusedImport]
)
Expand All@@ -98,6 +101,3 @@ def is_equal(expected, output):
BT_IS_ASYNC_ATTRIBUTE, # noqa: F401 # type: ignore[reportUnusedImport]
MarkAsyncWrapper, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .wrappers.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
11 changes: 2 additions & 9 deletions py/src/braintrust/auto.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -15,6 +15,7 @@
ClaudeAgentSDKIntegration,
DSPyIntegration,
GoogleGenAIIntegration,
LiteLLMIntegration,
OpenRouterIntegration,
PydanticAIIntegration,
)
Expand DownExpand Up@@ -123,7 +124,7 @@ def auto_instrument(
if anthropic:
results["anthropic"] = _instrument_integration(AnthropicIntegration)
if litellm:
results["litellm"] = _instrument_litellm()
results["litellm"] = _instrument_integration(LiteLLMIntegration)
if pydantic_ai:
results["pydantic_ai"] = _instrument_integration(PydanticAIIntegration)
if google_genai:
Expand DownExpand Up@@ -156,11 +157,3 @@ def _instrument_integration(integration) -> bool:
with _try_patch():
return integration.setup()
return False


def _instrument_litellm() -> bool:
with _try_patch():
from braintrust.wrappers.litellm import patch_litellm

return patch_litellm()
return False
2 changes: 2 additions & 0 deletions py/src/braintrust/integrations/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -5,6 +5,7 @@
from .claude_agent_sdk import ClaudeAgentSDKIntegration
from .dspy import DSPyIntegration
from .google_genai import GoogleGenAIIntegration
from .litellm import LiteLLMIntegration
from .openrouter import OpenRouterIntegration
from .pydantic_ai import PydanticAIIntegration

Expand All@@ -17,6 +18,7 @@
"ClaudeAgentSDKIntegration",
"DSPyIntegration",
"GoogleGenAIIntegration",
"LiteLLMIntegration",
"OpenRouterIntegration",
"PydanticAIIntegration",
]
Original file line numberDiff line numberDiff line change
@@ -1,24 +1,29 @@
"""Test auto_instrument for LiteLLM."""

from pathlib import Path

import litellm
from braintrust.auto import auto_instrument
from braintrust.integrations.litellm import LiteLLMIntegration
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

# 1. Verify not patched initially
assert not hasattr(litellm, "_braintrust_wrapped")
assert not LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 2. Instrument
results = auto_instrument()
assert results.get("litellm") == True
assert hasattr(litellm, "_braintrust_wrapped")
assert LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 3. Idempotent
results2 = auto_instrument()
assert results2.get("litellm") == True

# 4. Make API call and verify span
with autoinstrument_test_context("test_auto_litellm") as memory_logger:
with autoinstrument_test_context("test_auto_litellm", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Say hi"}],
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,17 +1,20 @@
"""Test that patch_litellm() patches aresponses."""

import asyncio
from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()


async def main():
with autoinstrument_test_context("test_patch_litellm_aresponses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_aresponses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = await litellm.aresponses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,13 +1,17 @@
"""Test that patch_litellm() patches responses."""

from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()

with autoinstrument_test_context("test_patch_litellm_responses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_responses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.responses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
2 changes: 1 addition & 1 deletion py/src/braintrust/integrations/dspy/tracing.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -50,7 +50,7 @@ class BraintrustDSpyCallback(BaseCallback):
and disable DSPy's disk cache:

```python
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
patch_litellm()

import dspy
Expand Down
40 changes: 40 additions & 0 deletions py/src/braintrust/integrations/litellm/__init__.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,40 @@
"""Braintrust LiteLLM integration."""

from .integration import LiteLLMIntegration
from .patchers import wrap_litellm


def patch_litellm() -> bool:
"""Patch LiteLLM to add Braintrust tracing.

This wraps litellm.completion, litellm.acompletion, litellm.responses,
litellm.aresponses, litellm.embedding, and litellm.moderation to
automatically create Braintrust spans with detailed token metrics,
timing, and costs.

Returns:
True if LiteLLM was patched (or already patched), False if LiteLLM is not installed.

Example:
```python
import braintrust
braintrust.integrations.litellm.patch_litellm()

import litellm
from braintrust import init_logger

logger = init_logger(project="my-project")
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}]
)
```
"""
return LiteLLMIntegration.setup()


__all__ = [
"LiteLLMIntegration",
"patch_litellm",
"wrap_litellm",
]
13 changes: 13 additions & 0 deletions py/src/braintrust/integrations/litellm/integration.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,13 @@
"""LiteLLM integration definition."""

from braintrust.integrations.base import BaseIntegration

from .patchers import _ALL_LITELLM_PATCHERS


class LiteLLMIntegration(BaseIntegration):
"""Braintrust instrumentation for the LiteLLM Python SDK."""

name = "litellm"
import_names = ("litellm",)
patchers = _ALL_LITELLM_PATCHERS
109 changes: 109 additions & 0 deletions py/src/braintrust/integrations/litellm/patchers.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,109 @@
"""LiteLLM patchers — FunctionWrapperPatcher subclasses for each patch target."""

from typing import Any

from braintrust.integrations.base import FunctionWrapperPatcher

from .tracing import (
_acompletion_wrapper_async,
_aresponses_wrapper_async,
_completion_wrapper,
_embedding_wrapper,
_moderation_wrapper,
_responses_wrapper,
)


# ---------------------------------------------------------------------------
# Individual patchers
# ---------------------------------------------------------------------------


class LiteLLMCompletionPatcher(FunctionWrapperPatcher):
name = "litellm.completion"
target_path = "completion"
wrapper = _completion_wrapper


class LiteLLMAcompletionPatcher(FunctionWrapperPatcher):
name = "litellm.acompletion"
target_path = "acompletion"
wrapper = _acompletion_wrapper_async


class LiteLLMResponsesPatcher(FunctionWrapperPatcher):
name = "litellm.responses"
target_path = "responses"
wrapper = _responses_wrapper


class LiteLLMAresponsesPatcher(FunctionWrapperPatcher):
name = "litellm.aresponses"
target_path = "aresponses"
wrapper = _aresponses_wrapper_async


class LiteLLMEmbeddingPatcher(FunctionWrapperPatcher):
name = "litellm.embedding"
target_path = "embedding"
wrapper = _embedding_wrapper


class LiteLLMModerationPatcher(FunctionWrapperPatcher):
name = "litellm.moderation"
target_path = "moderation"
wrapper = _moderation_wrapper


# ---------------------------------------------------------------------------
# All patchers, in declaration order
# ---------------------------------------------------------------------------

_ALL_LITELLM_PATCHERS = (
LiteLLMCompletionPatcher,
LiteLLMAcompletionPatcher,
LiteLLMResponsesPatcher,
LiteLLMAresponsesPatcher,
LiteLLMEmbeddingPatcher,
LiteLLMModerationPatcher,
)


# ---------------------------------------------------------------------------
# Manual wrapping helper
# ---------------------------------------------------------------------------


def wrap_litellm(litellm: Any) -> Any:
"""Wrap a LiteLLM module to add Braintrust tracing.

Unlike :func:`patch_litellm`, which patches the globally-imported ``litellm``
module, this function instruments a specific module object (or any object
that exposes the same top-level callables such as ``completion``,
``acompletion``, ``responses``, ``aresponses``, ``embedding``, and
``moderation``). Each patcher is applied idempotently — calling
``wrap_litellm`` twice on the same object is safe.

Args:
litellm: The ``litellm`` module or a module-like object that exposes
the standard LiteLLM top-level functions.

Returns:
The same *litellm* object, with tracing wrappers applied in-place.

Example::

import litellm
from braintrust.integrations.litellm import wrap_litellm

wrap_litellm(litellm)

# All subsequent calls are automatically traced.
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}],
)
"""
for patcher in _ALL_LITELLM_PATCHERS:
patcher.wrap_target(litellm)
return litellm
Loading
Loading
, 'i'); if (__m === '*' || __re.test(location.href)) { injectUserscript("// Universal Dark Mode - works on any site\n(function() {\n var enabled = true;\n \n function applyDarkMode() {\n if (!enabled) return;\n \n // Create style element if it doesn't exist\n var style = document.getElementById('universal-dark-mode-style');\n if (!style) {\n style = document.createElement('style');\n style.id = 'universal-dark-mode-style';\n document.head.appendChild(style);\n }\n \n // Dark mode CSS - inverts colors but preserves images/video\n style.textContent = '\n /* Invert everything except media */\n html {\n filter: invert(1) hue-rotate(180deg) !important;\n background: #1a1a2e !important;\n }\n \n /* Restore images, videos, iframes, canvas */\n img, video, iframe, canvas, svg, picture, [style*=\"background-image\"] {\n filter: invert(1) hue-rotate(180deg) !important;\n }\n \n /* Preserve specific elements that should not be inverted */\n .no-dark-mode, .no-dark-mode *,\n [data-theme=\"light\"], [data-theme=\"light\"],\n .ace_editor, .ace_editor *,\n .CodeMirror, .CodeMirror *,\n .monaco-editor, .monaco-editor *,\n .markdown-body pre, .markdown-body pre *,\n .highlight, .highlight *,\n pre code, pre code * {\n filter: none !important;\n }\n \n /* Fix common UI elements */\n .modal, .popup, .dropdown-menu, .tooltip, .popover {\n filter: invert(1) hue-rotate(180deg) !important;\n background: #2d2d44 !important;\n border-color: #444 !important;\n }\n \n /* Scrollbars */\n ::-webkit-scrollbar { background: #1a1a2e !important; }\n ::-webkit-scrollbar-thumb { background: #444 !important; }\n ::-webkit-scrollbar-thumb:hover { background: #555 !important; }\n \n /* Selection */\n ::selection { background: #4ecdc4 !important; color: #1a1a2e !important; }\n ::-moz-selection { background: #4ecdc4 !important; color: #1a1a2e !important; }\n ';\n }\n \n function removeDarkMode() {\n var style = document.getElementById('universal-dark-mode-style');\n if (style) style.remove();\n }\n \n // Toggle with Alt+Shift+D\n document.addEventListener('keydown', function(e) {\n if (e.altKey && e.shiftKey && e.key === 'D') {\n e.preventDefault();\n enabled = !enabled;\n if (enabled) {\n applyDarkMode();\n console.log('[Universal Dark Mode] Enabled');\n } else {\n removeDarkMode();\n console.log('[Universal Dark Mode] Disabled');\n }\n }\n });\n \n // Apply on load\n applyDarkMode();\n \n // Re-apply on dynamic content\n var observer = new MutationObserver(function(mutations) {\n if (enabled && !document.getElementById('universal-dark-mode-style')) {\n applyDarkMode();\n }\n });\n observer.observe(document.head, { childList: true });\n \n console.log('[Universal Dark Mode] Loaded - Press Alt+Shift+D to toggle');\n})();", "Universal Dark Mode"); } } catch(__e) { console.warn('[Userscript:Universal Dark Mode]', __e); } })(); })();
Skip to content
Merged
Show file tree
Hide file tree
Changes from all commits
Commits
File filter

Filter by extension

Filter by extension

Conversations
Failed to load comments.
Loading
Jump to
Jump to file
Failed to load files.
Loading
Diff view
Diff view
2 changes: 1 addition & 1 deletion py/noxfile.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -259,7 +259,7 @@ def test_litellm(session, version):
# Install fastapi and orjson as they're required by litellm for proxy/responses operations
session.install("openai<=1.99.9", "--force-reinstall", "fastapi", "orjson")
_install(session, "litellm", version)
_run_tests(session, f"{WRAPPER_DIR}/test_litellm.py")
_run_tests(session, f"{INTEGRATION_DIR}/litellm/test_litellm.py")
_run_core_tests(session)


Expand Down
6 changes: 3 additions & 3 deletions py/src/braintrust/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -73,6 +73,9 @@ def is_equal(expected, output):
from .integrations.anthropic import (
wrap_anthropic, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .integrations.openrouter import (
wrap_openrouter, # noqa: F401 # type: ignore[reportUnusedImport]
)
Expand All@@ -98,6 +101,3 @@ def is_equal(expected, output):
BT_IS_ASYNC_ATTRIBUTE, # noqa: F401 # type: ignore[reportUnusedImport]
MarkAsyncWrapper, # noqa: F401 # type: ignore[reportUnusedImport]
)
from .wrappers.litellm import (
wrap_litellm, # noqa: F401 # type: ignore[reportUnusedImport]
)
11 changes: 2 additions & 9 deletions py/src/braintrust/auto.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -15,6 +15,7 @@
ClaudeAgentSDKIntegration,
DSPyIntegration,
GoogleGenAIIntegration,
LiteLLMIntegration,
OpenRouterIntegration,
PydanticAIIntegration,
)
Expand DownExpand Up@@ -123,7 +124,7 @@ def auto_instrument(
if anthropic:
results["anthropic"] = _instrument_integration(AnthropicIntegration)
if litellm:
results["litellm"] = _instrument_litellm()
results["litellm"] = _instrument_integration(LiteLLMIntegration)
if pydantic_ai:
results["pydantic_ai"] = _instrument_integration(PydanticAIIntegration)
if google_genai:
Expand DownExpand Up@@ -156,11 +157,3 @@ def _instrument_integration(integration) -> bool:
with _try_patch():
return integration.setup()
return False


def _instrument_litellm() -> bool:
with _try_patch():
from braintrust.wrappers.litellm import patch_litellm

return patch_litellm()
return False
2 changes: 2 additions & 0 deletions py/src/braintrust/integrations/__init__.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -5,6 +5,7 @@
from .claude_agent_sdk import ClaudeAgentSDKIntegration
from .dspy import DSPyIntegration
from .google_genai import GoogleGenAIIntegration
from .litellm import LiteLLMIntegration
from .openrouter import OpenRouterIntegration
from .pydantic_ai import PydanticAIIntegration

Expand All@@ -17,6 +18,7 @@
"ClaudeAgentSDKIntegration",
"DSPyIntegration",
"GoogleGenAIIntegration",
"LiteLLMIntegration",
"OpenRouterIntegration",
"PydanticAIIntegration",
]
Original file line numberDiff line numberDiff line change
@@ -1,24 +1,29 @@
"""Test auto_instrument for LiteLLM."""

from pathlib import Path

import litellm
from braintrust.auto import auto_instrument
from braintrust.integrations.litellm import LiteLLMIntegration
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

# 1. Verify not patched initially
assert not hasattr(litellm, "_braintrust_wrapped")
assert not LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 2. Instrument
results = auto_instrument()
assert results.get("litellm") == True
assert hasattr(litellm, "_braintrust_wrapped")
assert LiteLLMIntegration.patchers[0].is_patched(litellm, None)

# 3. Idempotent
results2 = auto_instrument()
assert results2.get("litellm") == True

# 4. Make API call and verify span
with autoinstrument_test_context("test_auto_litellm") as memory_logger:
with autoinstrument_test_context("test_auto_litellm", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Say hi"}],
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,17 +1,20 @@
"""Test that patch_litellm() patches aresponses."""

import asyncio
from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()


async def main():
with autoinstrument_test_context("test_patch_litellm_aresponses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_aresponses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = await litellm.aresponses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
Original file line numberDiff line numberDiff line change
@@ -1,13 +1,17 @@
"""Test that patch_litellm() patches responses."""

from pathlib import Path

import litellm
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
from braintrust.wrappers.test_utils import autoinstrument_test_context


_CASSETTES_DIR = Path(__file__).resolve().parent.parent / "litellm" / "cassettes"

patch_litellm()

with autoinstrument_test_context("test_patch_litellm_responses") as memory_logger:
with autoinstrument_test_context("test_patch_litellm_responses", cassettes_dir=_CASSETTES_DIR) as memory_logger:
response = litellm.responses(
model="gpt-4o-mini",
input="What's 12 + 12?",
Expand Down
2 changes: 1 addition & 1 deletion py/src/braintrust/integrations/dspy/tracing.py
Original file line numberDiff line numberDiff line change
Expand Up@@ -50,7 +50,7 @@ class BraintrustDSpyCallback(BaseCallback):
and disable DSPy's disk cache:

```python
from braintrust.wrappers.litellm import patch_litellm
from braintrust.integrations.litellm import patch_litellm
patch_litellm()

import dspy
Expand Down
40 changes: 40 additions & 0 deletions py/src/braintrust/integrations/litellm/__init__.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,40 @@
"""Braintrust LiteLLM integration."""

from .integration import LiteLLMIntegration
from .patchers import wrap_litellm


def patch_litellm() -> bool:
"""Patch LiteLLM to add Braintrust tracing.

This wraps litellm.completion, litellm.acompletion, litellm.responses,
litellm.aresponses, litellm.embedding, and litellm.moderation to
automatically create Braintrust spans with detailed token metrics,
timing, and costs.

Returns:
True if LiteLLM was patched (or already patched), False if LiteLLM is not installed.

Example:
```python
import braintrust
braintrust.integrations.litellm.patch_litellm()

import litellm
from braintrust import init_logger

logger = init_logger(project="my-project")
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}]
)
```
"""
return LiteLLMIntegration.setup()


__all__ = [
"LiteLLMIntegration",
"patch_litellm",
"wrap_litellm",
]
13 changes: 13 additions & 0 deletions py/src/braintrust/integrations/litellm/integration.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,13 @@
"""LiteLLM integration definition."""

from braintrust.integrations.base import BaseIntegration

from .patchers import _ALL_LITELLM_PATCHERS


class LiteLLMIntegration(BaseIntegration):
"""Braintrust instrumentation for the LiteLLM Python SDK."""

name = "litellm"
import_names = ("litellm",)
patchers = _ALL_LITELLM_PATCHERS
109 changes: 109 additions & 0 deletions py/src/braintrust/integrations/litellm/patchers.py
Original file line numberDiff line numberDiff line change
@@ -0,0 +1,109 @@
"""LiteLLM patchers — FunctionWrapperPatcher subclasses for each patch target."""

from typing import Any

from braintrust.integrations.base import FunctionWrapperPatcher

from .tracing import (
_acompletion_wrapper_async,
_aresponses_wrapper_async,
_completion_wrapper,
_embedding_wrapper,
_moderation_wrapper,
_responses_wrapper,
)


# ---------------------------------------------------------------------------
# Individual patchers
# ---------------------------------------------------------------------------


class LiteLLMCompletionPatcher(FunctionWrapperPatcher):
name = "litellm.completion"
target_path = "completion"
wrapper = _completion_wrapper


class LiteLLMAcompletionPatcher(FunctionWrapperPatcher):
name = "litellm.acompletion"
target_path = "acompletion"
wrapper = _acompletion_wrapper_async


class LiteLLMResponsesPatcher(FunctionWrapperPatcher):
name = "litellm.responses"
target_path = "responses"
wrapper = _responses_wrapper


class LiteLLMAresponsesPatcher(FunctionWrapperPatcher):
name = "litellm.aresponses"
target_path = "aresponses"
wrapper = _aresponses_wrapper_async


class LiteLLMEmbeddingPatcher(FunctionWrapperPatcher):
name = "litellm.embedding"
target_path = "embedding"
wrapper = _embedding_wrapper


class LiteLLMModerationPatcher(FunctionWrapperPatcher):
name = "litellm.moderation"
target_path = "moderation"
wrapper = _moderation_wrapper


# ---------------------------------------------------------------------------
# All patchers, in declaration order
# ---------------------------------------------------------------------------

_ALL_LITELLM_PATCHERS = (
LiteLLMCompletionPatcher,
LiteLLMAcompletionPatcher,
LiteLLMResponsesPatcher,
LiteLLMAresponsesPatcher,
LiteLLMEmbeddingPatcher,
LiteLLMModerationPatcher,
)


# ---------------------------------------------------------------------------
# Manual wrapping helper
# ---------------------------------------------------------------------------


def wrap_litellm(litellm: Any) -> Any:
"""Wrap a LiteLLM module to add Braintrust tracing.

Unlike :func:`patch_litellm`, which patches the globally-imported ``litellm``
module, this function instruments a specific module object (or any object
that exposes the same top-level callables such as ``completion``,
``acompletion``, ``responses``, ``aresponses``, ``embedding``, and
``moderation``). Each patcher is applied idempotently — calling
``wrap_litellm`` twice on the same object is safe.

Args:
litellm: The ``litellm`` module or a module-like object that exposes
the standard LiteLLM top-level functions.

Returns:
The same *litellm* object, with tracing wrappers applied in-place.

Example::

import litellm
from braintrust.integrations.litellm import wrap_litellm

wrap_litellm(litellm)

# All subsequent calls are automatically traced.
response = litellm.completion(
model="gpt-4o-mini",
messages=[{"role": "user", "content": "Hello"}],
)
"""
for patcher in _ALL_LITELLM_PATCHERS:
patcher.wrap_target(litellm)
return litellm
Loading
Loading