all: support openai-compatible models The support is rather minimal at this point: Only hard-coded models, only -unsafe, only -skabandaddr="". The "shared" LLM package is strongly Claude-flavored. We can fix all of this and more over time, if we are inspired to. (Maybe we'll switch to https://github.com/maruel/genai?) The goal for now is to get the rough structure in place. I've rebased and rebuilt this more times than I care to remember.

commit: 4f84ab729ddbf428b0e891940f08f70b4edee05c [log] [tgz]
author: Josh Bleecher Snyder <josharian@gmail.com> Tue Apr 22 16:40:54 2025 -0700
committer: Josh Bleecher Snyder <josharian@gmail.com> Fri May 02 12:57:44 2025 -0700
tree: f2e52e4a01c188ada1f5acf8b2a013029b999495
parent: 44f9b4cec11e269a52fbfc099989ab425b8e125f [diff] [blame]
diff --git a/llm/ant/ant_test.go b/llm/ant/ant_test.go
new file mode 100644
index 0000000..67cc5db
--- /dev/null
+++ b/llm/ant/ant_test.go

@@ -0,0 +1,93 @@
+package ant
+
+import (
+	"math"
+	"testing"
+)
+
+// TestCalculateCostFromTokens tests the calculateCostFromTokens function
+func TestCalculateCostFromTokens(t *testing.T) {
+	tests := []struct {
+		name                     string
+		model                    string
+		inputTokens              uint64
+		outputTokens             uint64
+		cacheReadInputTokens     uint64
+		cacheCreationInputTokens uint64
+		want                     float64
+	}{
+		{
+			name:                     "Zero tokens",
+			model:                    Claude37Sonnet,
+			inputTokens:              0,
+			outputTokens:             0,
+			cacheReadInputTokens:     0,
+			cacheCreationInputTokens: 0,
+			want:                     0,
+		},
+		{
+			name:                     "1000 input tokens, 500 output tokens",
+			model:                    Claude37Sonnet,
+			inputTokens:              1000,
+			outputTokens:             500,
+			cacheReadInputTokens:     0,
+			cacheCreationInputTokens: 0,
+			want:                     0.0105,
+		},
+		{
+			name:                     "10000 input tokens, 5000 output tokens",
+			model:                    Claude37Sonnet,
+			inputTokens:              10000,
+			outputTokens:             5000,
+			cacheReadInputTokens:     0,
+			cacheCreationInputTokens: 0,
+			want:                     0.105,
+		},
+		{
+			name:                     "With cache read tokens",
+			model:                    Claude37Sonnet,
+			inputTokens:              1000,
+			outputTokens:             500,
+			cacheReadInputTokens:     2000,
+			cacheCreationInputTokens: 0,
+			want:                     0.0111,
+		},
+		{
+			name:                     "With cache creation tokens",
+			model:                    Claude37Sonnet,
+			inputTokens:              1000,
+			outputTokens:             500,
+			cacheReadInputTokens:     0,
+			cacheCreationInputTokens: 1500,
+			want:                     0.016125,
+		},
+		{
+			name:                     "With all token types",
+			model:                    Claude37Sonnet,
+			inputTokens:              1000,
+			outputTokens:             500,
+			cacheReadInputTokens:     2000,
+			cacheCreationInputTokens: 1500,
+			want:                     0.016725,
+		},
+	}
+
+	for _, tt := range tests {
+		t.Run(tt.name, func(t *testing.T) {
+			usage := usage{
+				InputTokens:              tt.inputTokens,
+				OutputTokens:             tt.outputTokens,
+				CacheReadInputTokens:     tt.cacheReadInputTokens,
+				CacheCreationInputTokens: tt.cacheCreationInputTokens,
+			}
+			mr := response{
+				Model: tt.model,
+				Usage: usage,
+			}
+			totalCost := mr.TotalDollars()
+			if math.Abs(totalCost-tt.want) > 0.0001 {
+				t.Errorf("totalCost = %v, want %v", totalCost, tt.want)
+			}
+		})
+	}
+}
commit	4f84ab729ddbf428b0e891940f08f70b4edee05c	[log] [tgz]
author	Josh Bleecher Snyder <josharian@gmail.com>	Tue Apr 22 16:40:54 2025 -0700
committer	Josh Bleecher Snyder <josharian@gmail.com>	Fri May 02 12:57:44 2025 -0700
tree	f2e52e4a01c188ada1f5acf8b2a013029b999495
parent	44f9b4cec11e269a52fbfc099989ab425b8e125f [diff] [blame]