| name | sweetviz-2-target-variable-analysis |
| description | Sub-skill of sweetviz: 2. Target Variable Analysis. |
| version | 1.0.0 |
| category | data-analysis |
| type | reference |
| scripts_exempt | true |
2. Target Variable Analysis
2. Target Variable Analysis
Binary Target Analysis:
import sweetviz as sv
import pandas as pd
import numpy as np
np.random.seed(42)
n = 3000
df = pd.DataFrame({
"feature_1": np.random.randn(n),
"feature_2": np.random.exponential(10, n),
"feature_3": np.random.choice(["A", "B", "C"], n),
"feature_4": np.random.randint(1, 100, n),
"target": np.random.choice([0, 1], n, p=[0.7, 0.3])
})
report = sv.analyze(
source=df,
target_feat="target"
)
report.show_html("target_analysis.html")
Continuous Target Analysis:
import sweetviz as sv
import pandas as pd
import numpy as np
np.random.seed(42)
n = 2000
x1 = np.random.randn(n)
x2 = np.random.exponential(5, n)
x3 = np.random.choice([0, 1], n)
target = 10 + 2*x1 + 0.5*x2 + 3*x3 + np.random.randn(n)
df = pd.DataFrame({
"feature_linear": x1,
"feature_exp": x2,
"feature_binary": x3,
"feature_noise": np.random.randn(n),
"price": target
})
report = sv.analyze(
source=df,
target_feat="price"
)
report.show_html("regression_target_analysis.html")
Multi-class Target:
import sweetviz as sv
import pandas as pd
import numpy as np
np.random.seed(42)
n = 2500
df = pd.DataFrame({
"feature_1": np.random.randn(n),
"feature_2": np.random.uniform(0, 100, n),
"category": np.random.choice(["A", "B", "C"], n),
"class_label": np.random.choice(
["Class_A", "Class_B", "Class_C", "Class_D"],
n, p=[0.4, 0.3, 0.2, 0.1]
)
})
report = sv.analyze(
source=df,
target_feat="class_label"
)
report.show_html("multiclass_analysis.html")