Code-Sleep-Python/Classification/code.py at master · mitansha/Code-Sleep-Python

88 lines (48 loc) · 1.98 KB

import pandas as pd
data = pd.read_csv('https://s3.amazonaws.com/demo-datasets/wine.csv')
df2 = data.drop('color', axis=1) # color is redundant
import numpy as np
numeric_data =  (numeric_data - np.mean(numeric_data))/(np.std(numeric_data))  
import sklearn.decomposition
pca = sklearn.decomposition.PCA(n_components=2)
principal_components = pca.fit(numeric_data).transform(numeric_data)
import matplotlib.pyplot as plt
from matplotlib.colors import ListedColormap
from matplotlib.backends.backend_pdf import PdfPages
observation_colormap = ListedColormap(['red', 'blue'])
x = principal_components[:,0]
y = principal_components[:,1]
plt.title("Principal Components of Wine")
plt.scatter(x, y, alpha = 0.2,
    c = data['high_quality'], cmap = observation_colormap, edgecolors = 'none')
plt.xlim(-8, 8); plt.ylim(-8, 8)
plt.xlabel("Principal Component 1"); plt.ylabel("Principal Component 2")
def accuracy(predictions, outcomes):
    # Enter your code here!
    occur =0
    for i in range(len(predictions)):
        if(predictions[i] == outcomes[i]):
            occur += 1
    return occur
x = np.array([1,2,3])
y = np.array([1,2,4])
print (accuracy(x,y))
print(accuracy(0,data["high_quality"]))
##############################   More accuracy
from sklearn.neighbors import KNeighborsClassifier
knn = KNeighborsClassifier(n_neighbors = 5)
knn.fit(numeric_data, data['high_quality'])
# Enter your code here!
library_predictions = knn.predict(numeric_data)
print(accuracy(library_predictions,data['high_quality']))
################
n_rows = data.shape[0]
random.seed(123)
selection = random.sample(range(n_rows), 10)
predictors = np.array(numeric_data)
training_indices = [i for i in range(len(predictors)) if i not in selection]
outcomes = np.array(data["high_quality"])
my_predictions = [knn_predict(p, predictors[training_indices,:], outcomes, k=5) for p in predictors[selection]]
percentage = accuracy(my_predictions,data.high_quality[selection])
print(percentage)

Provide feedback

Saved searches

Use saved searches to filter your results more quickly

FilesExpand file tree

code.py

Latest commit

History

code.py

File metadata and controls